Skip to content
This repository was archived by the owner on Apr 23, 2026. It is now read-only.

Commit 920f56b

Browse files
authored
Merge pull request #301 from verseles/update-arena-scores-1079658574268484239
Update LMArena English scores for Claude 4.6 and GPT 5.4 High
2 parents 5fc92de + d3766dd commit 920f56b

1 file changed

Lines changed: 119 additions & 4 deletions

File tree

data/showdown.json

Lines changed: 119 additions & 4 deletions
Original file line numberDiff line numberDiff line change
@@ -1,7 +1,7 @@
11
{
22
"meta": {
33
"version": "2026.04.16",
4-
"last_update": "2026-04-16T17:00:00Z",
4+
"last_update": "2026-04-18T19:00:00Z",
55
"schema_version": "1.0"
66
},
77
"models": [
@@ -119,6 +119,121 @@
119119
"arc_agi_2": 4.9
120120
}
121121
},
122+
{
123+
"id": "claude-opus-4-7-20260401-thinking-32k",
124+
"name": "Claude 4.7 Opus Thinking",
125+
"aka": ["claude-opus-4-7-thinking", "claude-4.7-opus-thinking", "claude-opus-4-7-max"],
126+
"superior_of": "claude-opus-4-7-20260401",
127+
"provider": "Anthropic",
128+
"type": "proprietary",
129+
"release_date": "2026-04-01",
130+
"pricing": {
131+
"input_per_1m": 5.0,
132+
"output_per_1m": 25.0,
133+
"average_per_1m": 10.0
134+
},
135+
"performance": {
136+
"output_speed_tps": 50,
137+
"latency_ttft_ms": 0,
138+
"source": "https://artificialanalysis.ai/models/claude-opus-4-7"
139+
},
140+
"editor_notes": "Thinking mode for Claude 4.7 Opus. Takes the #1 overall spot on LMArena.",
141+
"benchmark_scores": {
142+
"lmarena_en_elo": 1505,
143+
"lmarena_coding_elo": null,
144+
"lmarena_hard_elo": null,
145+
"lmarena_math_elo": null,
146+
"lmarena_creative_elo": null,
147+
"lmarena_if_elo": null,
148+
"swe_bench": null,
149+
"swe_bench_pro": null,
150+
"gpqa_diamond": null,
151+
"humanity_last_exam": null,
152+
"livebench": null,
153+
"math_500": null,
154+
"aime": null,
155+
"frontiermath": null,
156+
"arc_agi_2": null,
157+
"bfcl": null,
158+
"tau_bench": null,
159+
"osworld": null,
160+
"simpleqa": null,
161+
"mmmlu": null,
162+
"mmlu_pro": null,
163+
"live_code_bench": null,
164+
"terminal_bench": null,
165+
"webdev_arena_elo": null,
166+
"mathvista": null,
167+
"mmmu": null,
168+
"mmmu_pro": null,
169+
"lmarena_zh_elo": null,
170+
"lmarena_vision_elo": null,
171+
"livebench_reasoning": null,
172+
"livebench_coding": null,
173+
"livebench_agentic_coding": null,
174+
"livebench_math": null,
175+
"livebench_data_analysis": null,
176+
"livebench_language": null,
177+
"livebench_if": null
178+
}
179+
},
180+
{
181+
"id": "claude-opus-4-7-20260401",
182+
"name": "Claude 4.7 Opus",
183+
"aka": ["claude-opus-4-7", "claude-4.7-opus", "claude-4-7-opus"],
184+
"provider": "Anthropic",
185+
"type": "proprietary",
186+
"release_date": "2026-04-01",
187+
"pricing": {
188+
"input_per_1m": 5.0,
189+
"output_per_1m": 25.0,
190+
"average_per_1m": 10.0
191+
},
192+
"performance": {
193+
"output_speed_tps": 50,
194+
"latency_ttft_ms": 0,
195+
"source": "https://artificialanalysis.ai/models/claude-opus-4-7"
196+
},
197+
"editor_notes": "Anthropic's latest flagship model. Expensive and slower than average but boasts leading intelligence scores with a 1M token context window.",
198+
"benchmark_scores": {
199+
"lmarena_en_elo": 1498,
200+
"lmarena_coding_elo": null,
201+
"lmarena_hard_elo": null,
202+
"lmarena_math_elo": null,
203+
"lmarena_creative_elo": null,
204+
"lmarena_if_elo": null,
205+
"swe_bench": null,
206+
"swe_bench_pro": null,
207+
"gpqa_diamond": null,
208+
"humanity_last_exam": null,
209+
"livebench": null,
210+
"math_500": null,
211+
"aime": null,
212+
"frontiermath": null,
213+
"arc_agi_2": null,
214+
"bfcl": null,
215+
"tau_bench": null,
216+
"osworld": null,
217+
"simpleqa": null,
218+
"mmmlu": null,
219+
"mmlu_pro": null,
220+
"live_code_bench": null,
221+
"terminal_bench": null,
222+
"webdev_arena_elo": null,
223+
"mathvista": null,
224+
"mmmu": null,
225+
"mmmu_pro": null,
226+
"lmarena_zh_elo": null,
227+
"lmarena_vision_elo": null,
228+
"livebench_reasoning": null,
229+
"livebench_coding": null,
230+
"livebench_agentic_coding": null,
231+
"livebench_math": null,
232+
"livebench_data_analysis": null,
233+
"livebench_language": null,
234+
"livebench_if": null
235+
}
236+
},
122237
{
123238
"id": "claude-opus-4-6-20260205",
124239
"name": "Claude Opus 4.6",
@@ -149,7 +264,7 @@
149264
"livebench": null,
150265
"lmarena_coding_elo": 1547,
151266
"lmarena_creative_elo": 1468,
152-
"lmarena_en_elo": 1496,
267+
"lmarena_en_elo": 1497,
153268
"lmarena_hard_elo": 1529,
154269
"lmarena_if_elo": 1500,
155270
"lmarena_math_elo": 1501,
@@ -213,7 +328,7 @@
213328
"livebench": 76.33,
214329
"lmarena_coding_elo": 1556,
215330
"lmarena_creative_elo": 1493,
216-
"lmarena_en_elo": 1502,
331+
"lmarena_en_elo": 1503,
217332
"lmarena_hard_elo": 1536,
218333
"lmarena_if_elo": 1512,
219334
"lmarena_math_elo": 1512,
@@ -1053,7 +1168,7 @@
10531168
"livebench": 80.28,
10541169
"lmarena_coding_elo": 1532,
10551170
"lmarena_creative_elo": 1461,
1056-
"lmarena_en_elo": 1481,
1171+
"lmarena_en_elo": 1482,
10571172
"lmarena_hard_elo": 1507,
10581173
"lmarena_if_elo": 1488,
10591174
"lmarena_math_elo": 1522,

0 commit comments

Comments
 (0)