{"source":"BenchGecko","url":"https://benchgecko.ai/model/grok-3-mini","as_of":"2026-05-03","license":"BenchGecko collected data (prices, provider offers, Gecko Tests) is CC BY 4.0 · https://creativecommons.org/licenses/by/4.0/. Benchmark scores belong to their original publishers (see sources) and are aggregated with attribution. Attribution required: \"Source: BenchGecko\" with a link.","attribution":"Source: BenchGecko · https://benchgecko.ai/model/grok-3-mini","cite":"Grok-3 mini · benchmarks, pricing and providers. BenchGecko, data as of 2026-05-03. https://benchgecko.ai/model/grok-3-mini","slug":"grok-3-mini","name":"Grok-3 mini","provider":{"name":"xAI","slug":"xai"},"model_type":"text","release_date":"2024-01-01","is_open_source":false,"status":"benchmark-only","context_window":null,"description":null,"benchgecko_score":{"value":47.7,"rank":160,"of":312,"method":"Normalized average of public benchmark scores","method_url":"https://benchgecko.ai/methodology"},"avg_score":47.7,"list_price":{"input_usd_per_m":null,"output_usd_per_m":null,"source":"OpenRouter models API (checked daily)","as_of":"2026-10-05"},"pricing":{"input":null,"output":null},"providers":[],"cheapest_provider":null,"providers_source":null,"scores":[{"benchmark":"HELM — IFEval","benchmark_slug":"helm-ifeval","score":95.1,"unit":"%","category":"language","source":"Stanford HELM","source_url":"https://crfm.stanford.edu/helm/","benchmark_url":"https://benchgecko.ai/benchmark/helm-ifeval"},{"benchmark":"MATH level 5","benchmark_slug":"math-level-5","score":90.94,"unit":"%","category":"math","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/math-level-5"},{"benchmark":"HELM — MMLU-Pro","benchmark_slug":"helm-mmlu-pro","score":79.9,"unit":"%","category":"knowledge","source":"Stanford HELM","source_url":"https://crfm.stanford.edu/helm/","benchmark_url":"https://benchgecko.ai/benchmark/helm-mmlu-pro"},{"benchmark":"OTIS Mock AIME 2024-2025","benchmark_slug":"otis-mock-aime-2024-2025","score":77.76,"unit":"%","category":"math","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/otis-mock-aime-2024-2025"},{"benchmark":"Lech Mazur Writing","benchmark_slug":"lech-mazur-writing","score":73.5,"unit":"%","category":"knowledge","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/lech-mazur-writing"},{"benchmark":"GPQA diamond","benchmark_slug":"gpqa-diamond","score":68.35,"unit":"%","category":"knowledge","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/gpqa-diamond"},{"benchmark":"HELM — GPQA","benchmark_slug":"helm-gpqa","score":67.5,"unit":"%","category":"knowledge","source":"Stanford HELM","source_url":"https://crfm.stanford.edu/helm/","benchmark_url":"https://benchgecko.ai/benchmark/helm-gpqa"},{"benchmark":"Fiction.LiveBench","benchmark_slug":"fiction-livebench","score":66.7,"unit":"%","category":"knowledge","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/fiction-livebench"},{"benchmark":"HELM — WildBench","benchmark_slug":"helm-wildbench","score":65.1,"unit":"%","category":"reasoning","source":"Stanford HELM","source_url":"https://crfm.stanford.edu/helm/","benchmark_url":"https://benchgecko.ai/benchmark/helm-wildbench"},{"benchmark":"Aider polyglot","benchmark_slug":"aider-polyglot","score":49.3,"unit":"%","category":"coding","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/aider-polyglot"},{"benchmark":"WeirdML","benchmark_slug":"weirdml","score":42.58,"unit":"%","category":"coding","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/weirdml"},{"benchmark":"HELM — Omni-MATH","benchmark_slug":"helm-omni-math","score":31.8,"unit":"%","category":"math","source":"Stanford HELM","source_url":"https://crfm.stanford.edu/helm/","benchmark_url":"https://benchgecko.ai/benchmark/helm-omni-math"},{"benchmark":"SimpleQA Verified","benchmark_slug":"simpleqa-verified","score":21.1,"unit":"%","category":"knowledge","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/simpleqa-verified"},{"benchmark":"ARC-AGI","benchmark_slug":"arc-agi","score":16.5,"unit":"%","category":"reasoning","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/arc-agi"},{"benchmark":"FrontierMath-2025-02-28-Private","benchmark_slug":"frontiermath-2025-02-28-private","score":10.28,"unit":"%","category":"math","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/frontiermath-2025-02-28-private"},{"benchmark":"ARC-AGI-2","benchmark_slug":"arc-agi-2","score":0.42,"unit":"%","category":"reasoning","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/arc-agi-2"}],"gecko_tests":null,"links":{"page":"https://benchgecko.ai/model/grok-3-mini","json":"https://benchgecko.ai/api/v1/models/grok-3-mini","price_history":"https://benchgecko.ai/api/v1/price-history/grok-3-mini","provider_prices":null}}