{"source":"BenchGecko","url":"https://benchgecko.ai/model/gemini-2-5-pro","as_of":"2026-10-05","license":"BenchGecko collected data (prices, provider offers, Gecko Tests) is CC BY 4.0 · https://creativecommons.org/licenses/by/4.0/. Benchmark scores belong to their original publishers (see sources) and are aggregated with attribution. Attribution required: \"Source: BenchGecko\" with a link.","attribution":"Source: BenchGecko · https://benchgecko.ai/model/gemini-2-5-pro","cite":"Gemini 2.5 Pro · benchmarks, pricing and providers. BenchGecko, data as of 2026-10-05. https://benchgecko.ai/model/gemini-2-5-pro","slug":"gemini-2-5-pro","name":"Gemini 2.5 Pro","provider":{"name":"Google DeepMind","slug":"google"},"model_type":"multimodal","release_date":"2025-06-17","is_open_source":false,"status":"active","context_window":1048576,"description":"Gemini 2.5 Pro is Google’s state-of-the-art AI model designed for advanced reasoning, coding, mathematics, and scientific tasks. It employs “thinking” capabilities, enabling it to reason through responses with enhanced accuracy...","benchgecko_score":{"value":58.7,"rank":94,"of":312,"method":"Normalized average of public benchmark scores","method_url":"https://benchgecko.ai/methodology"},"avg_score":58.7,"list_price":{"input_usd_per_m":1.25,"output_usd_per_m":10,"source":"OpenRouter models API (checked daily)","as_of":"2026-10-05"},"pricing":{"input":1.25,"output":10},"providers":[{"provider":"Google AI Studio","endpoint":"google-ai-studio/flex","quantization":null,"context_tokens":1048576,"input_usd_per_m":0.625,"output_usd_per_m":5,"cache_read_usd_per_m":0.0625,"uptime_1d_pct":100,"as_of":"2026-10-05"},{"provider":"Google","endpoint":"google-vertex/eu","quantization":null,"context_tokens":1048576,"input_usd_per_m":1.25,"output_usd_per_m":10,"cache_read_usd_per_m":0.125,"uptime_1d_pct":93.67,"as_of":"2026-10-05"},{"provider":"Google","endpoint":"google-vertex/global","quantization":null,"context_tokens":1048576,"input_usd_per_m":1.25,"output_usd_per_m":10,"cache_read_usd_per_m":0.125,"uptime_1d_pct":98.65,"as_of":"2026-10-05"},{"provider":"Google","endpoint":"google-vertex/us","quantization":null,"context_tokens":1048576,"input_usd_per_m":1.25,"output_usd_per_m":10,"cache_read_usd_per_m":0.125,"uptime_1d_pct":96.14,"as_of":"2026-10-05"},{"provider":"Google AI Studio","endpoint":"google-ai-studio","quantization":null,"context_tokens":1048576,"input_usd_per_m":1.25,"output_usd_per_m":10,"cache_read_usd_per_m":0.125,"uptime_1d_pct":97.92,"as_of":"2026-10-05"},{"provider":"Google","endpoint":"google-vertex/global/priority","quantization":null,"context_tokens":1048576,"input_usd_per_m":2.25,"output_usd_per_m":18,"cache_read_usd_per_m":0.225,"uptime_1d_pct":99.88,"as_of":"2026-10-05"},{"provider":"Google AI Studio","endpoint":"google-ai-studio/priority","quantization":null,"context_tokens":1048576,"input_usd_per_m":2.25,"output_usd_per_m":18,"cache_read_usd_per_m":0.225,"uptime_1d_pct":100,"as_of":"2026-10-05"}],"cheapest_provider":{"provider":"Google AI Studio","endpoint":"google-ai-studio/flex","quantization":null,"context_tokens":1048576,"input_usd_per_m":0.625,"output_usd_per_m":5,"cache_read_usd_per_m":0.0625,"uptime_1d_pct":100,"as_of":"2026-10-05"},"providers_source":"OpenRouter endpoints API, one row per provider, refreshed daily","scores":[{"benchmark":"Chatbot Arena Elo — Overall","benchmark_slug":"arena-elo-overall","score":1445.57,"unit":"elo","category":"arena","source":"LMArena","source_url":"https://lmarena.ai/leaderboard","benchmark_url":"https://benchgecko.ai/benchmark/arena-elo-overall"},{"benchmark":"Chatbot Arena Elo — Coding","benchmark_slug":"arena-elo-coding","score":1226.98,"unit":"elo","category":"arena","source":"LMArena","source_url":"https://lmarena.ai/leaderboard","benchmark_url":"https://benchgecko.ai/benchmark/arena-elo-coding"},{"benchmark":"MATH level 5","benchmark_slug":"math-level-5","score":95.56,"unit":"%","category":"math","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/math-level-5"},{"benchmark":"Fiction.LiveBench","benchmark_slug":"fiction-livebench","score":91.7,"unit":"%","category":"knowledge","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/fiction-livebench"},{"benchmark":"OpenCompass — IFEval","benchmark_slug":"oc-ifeval","score":90,"unit":"%","category":"language","source":"OpenCompass","source_url":"https://rank.opencompass.org.cn","benchmark_url":"https://benchgecko.ai/benchmark/oc-ifeval"},{"benchmark":"OpenCompass — AIME2025","benchmark_slug":"oc-aime2025","score":88.7,"unit":"%","category":"math","source":"OpenCompass","source_url":"https://rank.opencompass.org.cn","benchmark_url":"https://benchgecko.ai/benchmark/oc-aime2025"},{"benchmark":"HELM — MMLU-Pro","benchmark_slug":"helm-mmlu-pro","score":86.3,"unit":"%","category":"knowledge","source":"Stanford HELM","source_url":"https://crfm.stanford.edu/helm/","benchmark_url":"https://benchgecko.ai/benchmark/helm-mmlu-pro"},{"benchmark":"OpenCompass — MMLU-Pro","benchmark_slug":"oc-mmlu-pro","score":85.8,"unit":"%","category":"knowledge","source":"OpenCompass","source_url":"https://rank.opencompass.org.cn","benchmark_url":"https://benchgecko.ai/benchmark/oc-mmlu-pro"},{"benchmark":"HELM — WildBench","benchmark_slug":"helm-wildbench","score":85.7,"unit":"%","category":"reasoning","source":"Stanford HELM","source_url":"https://crfm.stanford.edu/helm/","benchmark_url":"https://benchgecko.ai/benchmark/helm-wildbench"},{"benchmark":"OTIS Mock AIME 2024-2025","benchmark_slug":"otis-mock-aime-2024-2025","score":84.71,"unit":"%","category":"math","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/otis-mock-aime-2024-2025"},{"benchmark":"OpenCompass — GPQA-Diamond","benchmark_slug":"oc-gpqa-diamond","score":84.7,"unit":"%","category":"knowledge","source":"OpenCompass","source_url":"https://rank.opencompass.org.cn","benchmark_url":"https://benchgecko.ai/benchmark/oc-gpqa-diamond"},{"benchmark":"HELM — IFEval","benchmark_slug":"helm-ifeval","score":84,"unit":"%","category":"language","source":"Stanford HELM","source_url":"https://crfm.stanford.edu/helm/","benchmark_url":"https://benchgecko.ai/benchmark/helm-ifeval"},{"benchmark":"Lech Mazur Writing","benchmark_slug":"lech-mazur-writing","score":83.8,"unit":"%","category":"knowledge","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/lech-mazur-writing"},{"benchmark":"Aider polyglot","benchmark_slug":"aider-polyglot","score":83.1,"unit":"%","category":"coding","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/aider-polyglot"},{"benchmark":"GeoBench","benchmark_slug":"geobench","score":81,"unit":"%","category":"knowledge","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/geobench"},{"benchmark":"GPQA diamond","benchmark_slug":"gpqa-diamond","score":80.39,"unit":"%","category":"knowledge","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/gpqa-diamond"},{"benchmark":"HELM — GPQA","benchmark_slug":"helm-gpqa","score":74.9,"unit":"%","category":"knowledge","source":"Stanford HELM","source_url":"https://crfm.stanford.edu/helm/","benchmark_url":"https://benchgecko.ai/benchmark/helm-gpqa"},{"benchmark":"OpenCompass — LiveCodeBenchV6","benchmark_slug":"oc-livecodebenchv6","score":71.3,"unit":"%","category":"coding","source":"OpenCompass","source_url":"https://rank.opencompass.org.cn","benchmark_url":"https://benchgecko.ai/benchmark/oc-livecodebenchv6"},{"benchmark":"Dtbench","benchmark_slug":"dtbench","score":70.67,"unit":"%","category":"general","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/dtbench"},{"benchmark":"CadEval","benchmark_slug":"cadeval","score":64,"unit":"%","category":"coding","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/cadeval"},{"benchmark":"SWE-Bench verified","benchmark_slug":"swe-bench-verified","score":57.56,"unit":"%","category":"coding","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/swe-bench-verified"},{"benchmark":"SimpleQA Verified","benchmark_slug":"simpleqa-verified","score":56,"unit":"%","category":"knowledge","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/simpleqa-verified"},{"benchmark":"Metr Time Horizons","benchmark_slug":"metr-time-horizons","score":55.44,"unit":"%","category":"general","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/metr-time-horizons"},{"benchmark":"SimpleBench","benchmark_slug":"simplebench","score":54.88,"unit":"%","category":"reasoning","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/simplebench"},{"benchmark":"WeirdML","benchmark_slug":"weirdml","score":54.03,"unit":"%","category":"coding","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/weirdml"},{"benchmark":"AudioMultiChallenge","benchmark_slug":"seal-audiomultichallenge","score":46.9,"unit":"%","category":"knowledge","source":"Scale SEAL leaderboards","source_url":"https://scale.com/leaderboard","benchmark_url":"https://benchgecko.ai/benchmark/seal-audiomultichallenge"},{"benchmark":"AudioMultiChallenge — Text Output","benchmark_slug":"seal-audiomultichallenge-text-output","score":46.9,"unit":"%","category":"knowledge","source":"Scale SEAL leaderboards","source_url":"https://scale.com/leaderboard","benchmark_url":"https://benchgecko.ai/benchmark/seal-audiomultichallenge-text-output"},{"benchmark":"Balrog","benchmark_slug":"balrog","score":43.3,"unit":"%","category":"knowledge","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/balrog"},{"benchmark":"DeepResearch Bench","benchmark_slug":"deepresearch-bench","score":42.8,"unit":"%","category":"knowledge","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/deepresearch-bench"},{"benchmark":"HELM — Omni-MATH","benchmark_slug":"helm-omni-math","score":41.6,"unit":"%","category":"math","source":"Stanford HELM","source_url":"https://crfm.stanford.edu/helm/","benchmark_url":"https://benchgecko.ai/benchmark/helm-omni-math"},{"benchmark":"ARC-AGI","benchmark_slug":"arc-agi","score":41,"unit":"%","category":"reasoning","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/arc-agi"},{"benchmark":"Lmca","benchmark_slug":"lmca","score":40.91,"unit":"%","category":"general","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/lmca"},{"benchmark":"Artificial Analysis — Agentic Index","benchmark_slug":"aa-agentic-index","score":32.68,"unit":"index","category":"speed","source":"Artificial Analysis","source_url":"https://artificialanalysis.ai","benchmark_url":"https://benchgecko.ai/benchmark/aa-agentic-index"},{"benchmark":"Terminal Bench","benchmark_slug":"terminal-bench","score":32.64,"unit":"%","category":"coding","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/terminal-bench"},{"benchmark":"Artificial Analysis — Coding Index","benchmark_slug":"aa-coding-index","score":31.95,"unit":"index","category":"speed","source":"Artificial Analysis","source_url":"https://artificialanalysis.ai","benchmark_url":"https://benchgecko.ai/benchmark/aa-coding-index"},{"benchmark":"The Agent Company","benchmark_slug":"the-agent-company","score":30.3,"unit":"%","category":"agentic","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/the-agent-company"},{"benchmark":"Artificial Analysis — Quality Index","benchmark_slug":"aa-quality-index","score":26.98,"unit":"index","category":"speed","source":"Artificial Analysis","source_url":"https://artificialanalysis.ai","benchmark_url":"https://benchgecko.ai/benchmark/aa-quality-index"},{"benchmark":"FrontierMath-Tiers-1-3-v2-Private","benchmark_slug":"frontiermath-tiers-1-3-v2-private","score":24.56,"unit":"%","category":"math","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/frontiermath-tiers-1-3-v2-private"},{"benchmark":"Gdpval","benchmark_slug":"gdpval","score":23.3,"unit":"%","category":"general","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/gdpval"},{"benchmark":"OpenCompass — HLE","benchmark_slug":"oc-hle","score":21.1,"unit":"%","category":"knowledge","source":"OpenCompass","source_url":"https://rank.opencompass.org.cn","benchmark_url":"https://benchgecko.ai/benchmark/oc-hle"},{"benchmark":"VPCT","benchmark_slug":"vpct","score":19.6,"unit":"%","category":"knowledge","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/vpct"},{"benchmark":"HLE","benchmark_slug":"hle","score":17.69,"unit":"%","category":"knowledge","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/hle"},{"benchmark":"Chess Puzzles","benchmark_slug":"chess-puzzles","score":15.82,"unit":"%","category":"knowledge","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/chess-puzzles"},{"benchmark":"FrontierMath-2025-02-28-Private","benchmark_slug":"frontiermath-2025-02-28-private","score":14.14,"unit":"%","category":"math","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/frontiermath-2025-02-28-private"},{"benchmark":"APEX-Agents","benchmark_slug":"apex-agents","score":6.6,"unit":"%","category":"agentic","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/apex-agents"},{"benchmark":"ARC-AGI-2","benchmark_slug":"arc-agi-2","score":4.86,"unit":"%","category":"reasoning","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/arc-agi-2"},{"benchmark":"FrontierMath-Tier-4-2025-07-01-Private","benchmark_slug":"frontiermath-tier-4-2025-07-01-private","score":4.17,"unit":"%","category":"math","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/frontiermath-tier-4-2025-07-01-private"},{"benchmark":"GSO-Bench","benchmark_slug":"gso-bench","score":3.92,"unit":"%","category":"coding","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/gso-bench"},{"benchmark":"Remote Labor Index","benchmark_slug":"remote-labor-index","score":0.83,"unit":"%","category":"general","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/remote-labor-index"},{"benchmark":"FrontierMath-Tier-4-v2-Private","benchmark_slug":"frontiermath-tier-4-v2-private","score":0,"unit":"%","category":"math","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/frontiermath-tier-4-v2-private"}],"gecko_tests":null,"links":{"page":"https://benchgecko.ai/model/gemini-2-5-pro","json":"https://benchgecko.ai/api/v1/models/gemini-2-5-pro","price_history":"https://benchgecko.ai/api/v1/price-history/gemini-2-5-pro","provider_prices":"https://benchgecko.ai/pricing/arbitrage/gemini-2-5-pro"}}