{"source":"BenchGecko","url":"https://benchgecko.ai/model/claude-opus-4-8","as_of":"2026-10-05","license":"BenchGecko collected data (prices, provider offers, Gecko Tests) is CC BY 4.0 · https://creativecommons.org/licenses/by/4.0/. Benchmark scores belong to their original publishers (see sources) and are aggregated with attribution. Attribution required: \"Source: BenchGecko\" with a link.","attribution":"Source: BenchGecko · https://benchgecko.ai/model/claude-opus-4-8","cite":"Claude Opus 4.8 · benchmarks, pricing and providers. BenchGecko, data as of 2026-10-05. https://benchgecko.ai/model/claude-opus-4-8","slug":"claude-opus-4-8","name":"Claude Opus 4.8","provider":{"name":"Anthropic","slug":"anthropic"},"model_type":"multimodal","release_date":"2026-05-27","is_open_source":false,"status":"active","context_window":1000000,"description":"Claude Opus 4.8 is Anthropic's most capable generally available model in the Opus family. It supports text, image, and file inputs with text output, with reasoning support and a 1M-token...","benchgecko_score":{"value":71.8,"rank":33,"of":312,"method":"Normalized average of public benchmark scores","method_url":"https://benchgecko.ai/methodology"},"avg_score":71.8,"list_price":{"input_usd_per_m":5,"output_usd_per_m":25,"source":"OpenRouter models API (checked daily)","as_of":"2026-10-05"},"pricing":{"input":5,"output":25},"providers":[],"cheapest_provider":null,"providers_source":null,"scores":[{"benchmark":"Chatbot Arena Elo — Coding","benchmark_slug":"arena-elo-coding","score":1534.13,"unit":"elo","category":"arena","source":"LMArena","source_url":"https://lmarena.ai/leaderboard","benchmark_url":"https://benchgecko.ai/benchmark/arena-elo-coding"},{"benchmark":"Chatbot Arena Elo — Overall","benchmark_slug":"arena-elo-overall","score":1474.54,"unit":"elo","category":"arena","source":"LMArena","source_url":"https://lmarena.ai/leaderboard","benchmark_url":"https://benchgecko.ai/benchmark/arena-elo-overall"},{"benchmark":"OTIS Mock AIME 2024-2025","benchmark_slug":"otis-mock-aime-2024-2025","score":98.33,"unit":"%","category":"math","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/otis-mock-aime-2024-2025"},{"benchmark":"ARC-AGI","benchmark_slug":"arc-agi","score":92.5,"unit":"%","category":"reasoning","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/arc-agi"},{"benchmark":"Dtbench","benchmark_slug":"dtbench","score":91.55,"unit":"%","category":"general","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/dtbench"},{"benchmark":"GPQA diamond","benchmark_slug":"gpqa-diamond","score":88.05,"unit":"%","category":"knowledge","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/gpqa-diamond"},{"benchmark":"Surface Evolver Bench","benchmark_slug":"surface-evolver-bench","score":87.5,"unit":"%","category":"general","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/surface-evolver-bench"},{"benchmark":"WeirdML","benchmark_slug":"weirdml","score":82.89,"unit":"%","category":"coding","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/weirdml"},{"benchmark":"FrontierMath-Tiers-1-3-v2-Private","benchmark_slug":"frontiermath-tiers-1-3-v2-private","score":80,"unit":"%","category":"math","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/frontiermath-tiers-1-3-v2-private"},{"benchmark":"ARC-AGI-2","benchmark_slug":"arc-agi-2","score":72.08,"unit":"%","category":"reasoning","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/arc-agi-2"},{"benchmark":"Proofbench","benchmark_slug":"proofbench","score":69,"unit":"%","category":"general","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/proofbench"},{"benchmark":"Lmca","benchmark_slug":"lmca","score":67.66,"unit":"%","category":"general","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/lmca"},{"benchmark":"Deepswe","benchmark_slug":"deepswe","score":58.97,"unit":"%","category":"coding","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/deepswe"},{"benchmark":"SimpleBench","benchmark_slug":"simplebench","score":57.76,"unit":"%","category":"reasoning","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/simplebench"},{"benchmark":"FrontierMath-Tier-4-v2-Private","benchmark_slug":"frontiermath-tier-4-v2-private","score":56.1,"unit":"%","category":"math","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/frontiermath-tier-4-v2-private"},{"benchmark":"SimpleQA Verified","benchmark_slug":"simpleqa-verified","score":53,"unit":"%","category":"knowledge","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/simpleqa-verified"},{"benchmark":"DeepResearch Bench","benchmark_slug":"deepresearch-bench","score":50.2,"unit":"%","category":"knowledge","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/deepresearch-bench"},{"benchmark":"APEX-Agents","benchmark_slug":"apex-agents","score":48.9,"unit":"%","category":"agentic","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/apex-agents"},{"benchmark":"FrontierMath-2025-02-28-Private","benchmark_slug":"frontiermath-2025-02-28-private","score":47.24,"unit":"%","category":"math","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/frontiermath-2025-02-28-private"},{"benchmark":"GSO-Bench","benchmark_slug":"gso-bench","score":47.06,"unit":"%","category":"coding","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/gso-bench"},{"benchmark":"Frontiercode","benchmark_slug":"frontiercode","score":46.5,"unit":"%","category":"coding","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/frontiercode"},{"benchmark":"PostTrainBench","benchmark_slug":"posttrainbench","score":33.84,"unit":"%","category":"knowledge","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/posttrainbench"},{"benchmark":"FrontierMath-Tier-4-2025-07-01-Private","benchmark_slug":"frontiermath-tier-4-2025-07-01-private","score":31.25,"unit":"%","category":"math","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/frontiermath-tier-4-2025-07-01-private"},{"benchmark":"Chess Puzzles","benchmark_slug":"chess-puzzles","score":30.56,"unit":"%","category":"knowledge","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/chess-puzzles"},{"benchmark":"Mystery Game Puzzles","benchmark_slug":"mystery-game-puzzles","score":29.5,"unit":"%","category":"general","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/mystery-game-puzzles"},{"benchmark":"Ebr Bench","benchmark_slug":"ebr-bench","score":28.57,"unit":"%","category":"general","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/ebr-bench"},{"benchmark":"Osworld 2 0","benchmark_slug":"osworld-2-0","score":20.6,"unit":"%","category":"general","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/osworld-2-0"},{"benchmark":"Furniture Assembly","benchmark_slug":"furniture-assembly","score":17.86,"unit":"%","category":"general","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/furniture-assembly"},{"benchmark":"Remote Labor Index","benchmark_slug":"remote-labor-index","score":8.33,"unit":"%","category":"general","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/remote-labor-index"}],"gecko_tests":null,"links":{"page":"https://benchgecko.ai/model/claude-opus-4-8","json":"https://benchgecko.ai/api/v1/models/claude-opus-4-8","price_history":"https://benchgecko.ai/api/v1/price-history/claude-opus-4-8","provider_prices":null}}