{"source":"BenchGecko","url":"https://benchgecko.ai/model/gemini-3-1-pro-preview","as_of":"2026-10-05","license":"BenchGecko collected data (prices, provider offers, Gecko Tests) is CC BY 4.0 · https://creativecommons.org/licenses/by/4.0/. Benchmark scores belong to their original publishers (see sources) and are aggregated with attribution. Attribution required: \"Source: BenchGecko\" with a link.","attribution":"Source: BenchGecko · https://benchgecko.ai/model/gemini-3-1-pro-preview","cite":"Gemini 3.1 Pro Preview · benchmarks, pricing and providers. BenchGecko, data as of 2026-10-05. https://benchgecko.ai/model/gemini-3-1-pro-preview","slug":"gemini-3-1-pro-preview","name":"Gemini 3.1 Pro Preview","provider":{"name":"Google DeepMind","slug":"google"},"model_type":"multimodal","release_date":"2026-02-19","is_open_source":false,"status":"preview","context_window":1048576,"description":"Gemini 3.1 Pro Preview is Google’s frontier reasoning model, delivering enhanced software engineering performance, improved agentic reliability, and more efficient token usage across complex workflows. Building on the multimodal foundation...","benchgecko_score":{"value":65,"rank":57,"of":312,"method":"Normalized average of public benchmark scores","method_url":"https://benchgecko.ai/methodology"},"avg_score":65,"list_price":{"input_usd_per_m":2,"output_usd_per_m":12,"source":"OpenRouter models API (checked daily)","as_of":"2026-10-05"},"pricing":{"input":2,"output":12},"providers":[{"provider":"Google","endpoint":"google-vertex/global/flex","quantization":null,"context_tokens":1048576,"input_usd_per_m":1,"output_usd_per_m":6,"cache_read_usd_per_m":0.1,"uptime_1d_pct":95.1,"as_of":"2026-10-05"},{"provider":"Google AI Studio","endpoint":"google-ai-studio/flex","quantization":null,"context_tokens":1048576,"input_usd_per_m":1,"output_usd_per_m":6,"cache_read_usd_per_m":0.1,"uptime_1d_pct":99.99,"as_of":"2026-10-05"},{"provider":"Google","endpoint":"google-vertex/global","quantization":null,"context_tokens":1048576,"input_usd_per_m":2,"output_usd_per_m":12,"cache_read_usd_per_m":0.2,"uptime_1d_pct":99.61,"as_of":"2026-10-05"},{"provider":"Google AI Studio","endpoint":"google-ai-studio","quantization":null,"context_tokens":1048576,"input_usd_per_m":2,"output_usd_per_m":12,"cache_read_usd_per_m":0.2,"uptime_1d_pct":99.75,"as_of":"2026-10-05"},{"provider":"Google","endpoint":"google-vertex/global/priority","quantization":null,"context_tokens":1048576,"input_usd_per_m":3.6,"output_usd_per_m":21.6,"cache_read_usd_per_m":0.36,"uptime_1d_pct":null,"as_of":"2026-10-05"},{"provider":"Google AI Studio","endpoint":"google-ai-studio/priority","quantization":null,"context_tokens":1048576,"input_usd_per_m":3.6,"output_usd_per_m":21.6,"cache_read_usd_per_m":0.36,"uptime_1d_pct":99.75,"as_of":"2026-10-05"}],"cheapest_provider":{"provider":"Google","endpoint":"google-vertex/global/flex","quantization":null,"context_tokens":1048576,"input_usd_per_m":1,"output_usd_per_m":6,"cache_read_usd_per_m":0.1,"uptime_1d_pct":95.1,"as_of":"2026-10-05"},"providers_source":"OpenRouter endpoints API, one row per provider, refreshed daily","scores":[{"benchmark":"Chatbot Arena Elo — Overall","benchmark_slug":"arena-elo-overall","score":1486.98,"unit":"elo","category":"arena","source":"LMArena","source_url":"https://lmarena.ai/leaderboard","benchmark_url":"https://benchgecko.ai/benchmark/arena-elo-overall"},{"benchmark":"Chatbot Arena Elo — Coding","benchmark_slug":"arena-elo-coding","score":1446.2,"unit":"elo","category":"arena","source":"LMArena","source_url":"https://lmarena.ai/leaderboard","benchmark_url":"https://benchgecko.ai/benchmark/arena-elo-coding"},{"benchmark":"ARC-AGI","benchmark_slug":"arc-agi","score":98,"unit":"%","category":"reasoning","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/arc-agi"},{"benchmark":"OTIS Mock AIME 2024-2025","benchmark_slug":"otis-mock-aime-2024-2025","score":95.6,"unit":"%","category":"math","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/otis-mock-aime-2024-2025"},{"benchmark":"Artificial Analysis · tau2-Bench Telecom","benchmark_slug":"aa-tau2-bench","score":95.6,"unit":"%","category":"speed","source":"Artificial Analysis","source_url":"https://artificialanalysis.ai","benchmark_url":"https://benchgecko.ai/benchmark/aa-tau2-bench"},{"benchmark":"Dtbench","benchmark_slug":"dtbench","score":95.12,"unit":"%","category":"general","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/dtbench"},{"benchmark":"Artificial Analysis · GPQA Diamond","benchmark_slug":"aa-gpqa-diamond","score":94.1,"unit":"%","category":"speed","source":"Artificial Analysis","source_url":"https://artificialanalysis.ai","benchmark_url":"https://benchgecko.ai/benchmark/aa-gpqa-diamond"},{"benchmark":"GPQA diamond","benchmark_slug":"gpqa-diamond","score":92.59,"unit":"%","category":"knowledge","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/gpqa-diamond"},{"benchmark":"Artificial Analysis · MMMU Pro","benchmark_slug":"aa-mmmu-pro","score":82.4,"unit":"%","category":"speed","source":"Artificial Analysis","source_url":"https://artificialanalysis.ai","benchmark_url":"https://benchgecko.ai/benchmark/aa-mmmu-pro"},{"benchmark":"Artificial Analysis · Long Context Reasoning","benchmark_slug":"aa-long-context-reasoning","score":82,"unit":"%","category":"speed","source":"Artificial Analysis","source_url":"https://artificialanalysis.ai","benchmark_url":"https://benchgecko.ai/benchmark/aa-long-context-reasoning"},{"benchmark":"Terminal Bench","benchmark_slug":"terminal-bench","score":80.22,"unit":"%","category":"coding","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/terminal-bench"},{"benchmark":"ARC-AGI-2","benchmark_slug":"arc-agi-2","score":77.1,"unit":"%","category":"reasoning","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/arc-agi-2"},{"benchmark":"Artificial Analysis · IFBench","benchmark_slug":"aa-ifbench","score":77.1,"unit":"%","category":"speed","source":"Artificial Analysis","source_url":"https://artificialanalysis.ai","benchmark_url":"https://benchgecko.ai/benchmark/aa-ifbench"},{"benchmark":"Metr Time Horizons","benchmark_slug":"metr-time-horizons","score":77.04,"unit":"%","category":"general","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/metr-time-horizons"},{"benchmark":"SWE-Bench verified","benchmark_slug":"swe-bench-verified","score":75.62,"unit":"%","category":"coding","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/swe-bench-verified"},{"benchmark":"SimpleBench","benchmark_slug":"simplebench","score":75.52,"unit":"%","category":"reasoning","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/simplebench"},{"benchmark":"SimpleQA Verified","benchmark_slug":"simpleqa-verified","score":73.5,"unit":"%","category":"knowledge","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/simpleqa-verified"},{"benchmark":"WeirdML","benchmark_slug":"weirdml","score":72.07,"unit":"%","category":"coding","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/weirdml"},{"benchmark":"MultiChallenge","benchmark_slug":"seal-multichallenge","score":71.37,"unit":"%","category":"knowledge","source":"Scale SEAL leaderboards","source_url":"https://scale.com/leaderboard","benchmark_url":"https://benchgecko.ai/benchmark/seal-multichallenge"},{"benchmark":"Artificial Analysis — Coding Index","benchmark_slug":"aa-coding-index","score":68.83,"unit":"index","category":"speed","source":"Artificial Analysis","source_url":"https://artificialanalysis.ai","benchmark_url":"https://benchgecko.ai/benchmark/aa-coding-index"},{"benchmark":"MultiNRC","benchmark_slug":"seal-multinrc","score":64.74,"unit":"%","category":"knowledge","source":"Scale SEAL leaderboards","source_url":"https://scale.com/leaderboard","benchmark_url":"https://benchgecko.ai/benchmark/seal-multinrc"},{"benchmark":"Lmca","benchmark_slug":"lmca","score":63.32,"unit":"%","category":"general","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/lmca"},{"benchmark":"FrontierMath-Tiers-1-3-v2-Private","benchmark_slug":"frontiermath-tiers-1-3-v2-private","score":59.65,"unit":"%","category":"math","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/frontiermath-tiers-1-3-v2-private"},{"benchmark":"Artificial Analysis · SciCode","benchmark_slug":"aa-scicode","score":58.7,"unit":"%","category":"speed","source":"Artificial Analysis","source_url":"https://artificialanalysis.ai","benchmark_url":"https://benchgecko.ai/benchmark/aa-scicode"},{"benchmark":"Balrog","benchmark_slug":"balrog","score":57,"unit":"%","category":"knowledge","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/balrog"},{"benchmark":"Artificial Analysis · Terminal-Bench Hard","benchmark_slug":"aa-terminal-bench-hard","score":53.8,"unit":"%","category":"speed","source":"Artificial Analysis","source_url":"https://artificialanalysis.ai","benchmark_url":"https://benchgecko.ai/benchmark/aa-terminal-bench-hard"},{"benchmark":"Chess Puzzles","benchmark_slug":"chess-puzzles","score":52.65,"unit":"%","category":"knowledge","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/chess-puzzles"},{"benchmark":"DeepResearch Bench","benchmark_slug":"deepresearch-bench","score":47.8,"unit":"%","category":"knowledge","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/deepresearch-bench"},{"benchmark":"Artificial Analysis · Humanity's Last Exam","benchmark_slug":"aa-humanitys-last-exam","score":47,"unit":"%","category":"speed","source":"Artificial Analysis","source_url":"https://artificialanalysis.ai","benchmark_url":"https://benchgecko.ai/benchmark/aa-humanitys-last-exam"},{"benchmark":"HLE","benchmark_slug":"hle","score":43.74,"unit":"%","category":"knowledge","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/hle"},{"benchmark":"FrontierMath-2025-02-28-Private","benchmark_slug":"frontiermath-2025-02-28-private","score":36.9,"unit":"%","category":"math","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/frontiermath-2025-02-28-private"},{"benchmark":"APEX-Agents","benchmark_slug":"apex-agents","score":35.3,"unit":"%","category":"agentic","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/apex-agents"},{"benchmark":"Artificial Analysis — Quality Index","benchmark_slug":"aa-quality-index","score":29.72,"unit":"index","category":"speed","source":"Artificial Analysis","source_url":"https://artificialanalysis.ai","benchmark_url":"https://benchgecko.ai/benchmark/aa-quality-index"},{"benchmark":"VisualToolBench (VTB)","benchmark_slug":"seal-visual-tool-bench","score":28.97,"unit":"%","category":"knowledge","source":"Scale SEAL leaderboards","source_url":"https://scale.com/leaderboard","benchmark_url":"https://benchgecko.ai/benchmark/seal-visual-tool-bench"},{"benchmark":"Mystery Game Puzzles","benchmark_slug":"mystery-game-puzzles","score":27.3,"unit":"%","category":"general","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/mystery-game-puzzles"},{"benchmark":"FrontierMath-Tier-4-v2-Private","benchmark_slug":"frontiermath-tier-4-v2-private","score":26.83,"unit":"%","category":"math","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/frontiermath-tier-4-v2-private"},{"benchmark":"Exploitbench","benchmark_slug":"exploitbench","score":26.1,"unit":"%","category":"general","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/exploitbench"},{"benchmark":"Proofbench","benchmark_slug":"proofbench","score":26,"unit":"%","category":"general","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/proofbench"},{"benchmark":"GSO-Bench","benchmark_slug":"gso-bench","score":22.55,"unit":"%","category":"coding","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/gso-bench"},{"benchmark":"PostTrainBench","benchmark_slug":"posttrainbench","score":21.99,"unit":"%","category":"knowledge","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/posttrainbench"},{"benchmark":"Artificial Analysis — Agentic Index","benchmark_slug":"aa-agentic-index","score":21.4,"unit":"index","category":"speed","source":"Artificial Analysis","source_url":"https://artificialanalysis.ai","benchmark_url":"https://benchgecko.ai/benchmark/aa-agentic-index"},{"benchmark":"Cl Bench","benchmark_slug":"cl-bench","score":20.8,"unit":"%","category":"general","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/cl-bench"},{"benchmark":"EnigmaEval","benchmark_slug":"seal-enigmaeval","score":19.76,"unit":"%","category":"knowledge","source":"Scale SEAL leaderboards","source_url":"https://scale.com/leaderboard","benchmark_url":"https://benchgecko.ai/benchmark/seal-enigmaeval"},{"benchmark":"Artificial Analysis · CritPt","benchmark_slug":"aa-critpt","score":17.7,"unit":"%","category":"speed","source":"Artificial Analysis","source_url":"https://artificialanalysis.ai","benchmark_url":"https://benchgecko.ai/benchmark/aa-critpt"},{"benchmark":"Cl Bench Life","benchmark_slug":"cl-bench-life","score":16.9,"unit":"%","category":"general","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/cl-bench-life"},{"benchmark":"FrontierMath-Tier-4-2025-07-01-Private","benchmark_slug":"frontiermath-tier-4-2025-07-01-private","score":16.7,"unit":"%","category":"math","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/frontiermath-tier-4-2025-07-01-private"},{"benchmark":"Ebr Bench","benchmark_slug":"ebr-bench","score":14.29,"unit":"%","category":"general","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/ebr-bench"},{"benchmark":"Artificial Analysis · GDPval","benchmark_slug":"aa-gdpval","score":13.8,"unit":"%","category":"speed","source":"Artificial Analysis","source_url":"https://artificialanalysis.ai","benchmark_url":"https://benchgecko.ai/benchmark/aa-gdpval"},{"benchmark":"Deepswe","benchmark_slug":"deepswe","score":11.73,"unit":"%","category":"coding","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/deepswe"},{"benchmark":"Mirrorcode","benchmark_slug":"mirrorcode","score":8.89,"unit":"%","category":"coding","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/mirrorcode"},{"benchmark":"Furniture Assembly","benchmark_slug":"furniture-assembly","score":0,"unit":"%","category":"general","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/furniture-assembly"}],"gecko_tests":null,"links":{"page":"https://benchgecko.ai/model/gemini-3-1-pro-preview","json":"https://benchgecko.ai/api/v1/models/gemini-3-1-pro-preview","price_history":"https://benchgecko.ai/api/v1/price-history/gemini-3-1-pro-preview","provider_prices":"https://benchgecko.ai/pricing/arbitrage/gemini-3-1-pro-preview"}}