{"source":"BenchGecko","url":"https://benchgecko.ai/model/qwen3-next-80b-a3b-thinking","as_of":"2026-10-05","license":"BenchGecko collected data (prices, provider offers, Gecko Tests) is CC BY 4.0 · https://creativecommons.org/licenses/by/4.0/. Benchmark scores belong to their original publishers (see sources) and are aggregated with attribution. Attribution required: \"Source: BenchGecko\" with a link.","attribution":"Source: BenchGecko · https://benchgecko.ai/model/qwen3-next-80b-a3b-thinking","cite":"Qwen3 Next 80B A3B Thinking · benchmarks, pricing and providers. BenchGecko, data as of 2026-10-05. https://benchgecko.ai/model/qwen3-next-80b-a3b-thinking","slug":"qwen3-next-80b-a3b-thinking","name":"Qwen3 Next 80B A3B Thinking","provider":{"name":"Alibaba Qwen","slug":"alibaba"},"model_type":"text","release_date":"2025-09-11","is_open_source":true,"status":"active","context_window":262144,"description":"Qwen3-Next-80B-A3B-Thinking is a reasoning-first chat model in the Qwen3-Next line that outputs structured “thinking” traces by default. It’s designed for hard multi-step problems; math proofs, code synthesis/debugging, logic, and agentic...","benchgecko_score":{"value":56.8,"rank":107,"of":312,"method":"Normalized average of public benchmark scores","method_url":"https://benchgecko.ai/methodology"},"avg_score":56.8,"list_price":{"input_usd_per_m":0.15,"output_usd_per_m":1.2,"source":"OpenRouter models API (checked daily)","as_of":"2026-10-05"},"pricing":{"input":0.15,"output":1.2},"providers":[{"provider":"Alibaba","endpoint":"alibaba","quantization":null,"context_tokens":131072,"input_usd_per_m":0.15,"output_usd_per_m":1.2,"cache_read_usd_per_m":null,"uptime_1d_pct":95.09,"as_of":"2026-10-05"},{"provider":"Google","endpoint":"google-vertex/global","quantization":null,"context_tokens":262144,"input_usd_per_m":0.15,"output_usd_per_m":1.2,"cache_read_usd_per_m":null,"uptime_1d_pct":99.76,"as_of":"2026-10-05"}],"cheapest_provider":{"provider":"Alibaba","endpoint":"alibaba","quantization":null,"context_tokens":131072,"input_usd_per_m":0.15,"output_usd_per_m":1.2,"cache_read_usd_per_m":null,"uptime_1d_pct":95.09,"as_of":"2026-10-05"},"providers_source":"OpenRouter endpoints API, one row per provider, refreshed daily","scores":[{"benchmark":"Chatbot Arena Elo — Overall","benchmark_slug":"arena-elo-overall","score":1369.35,"unit":"elo","category":"arena","source":"LMArena","source_url":"https://lmarena.ai/leaderboard","benchmark_url":"https://benchgecko.ai/benchmark/arena-elo-overall"},{"benchmark":"OpenCompass — IFEval","benchmark_slug":"oc-ifeval","score":89.5,"unit":"%","category":"language","source":"OpenCompass","source_url":"https://rank.opencompass.org.cn","benchmark_url":"https://benchgecko.ai/benchmark/oc-ifeval"},{"benchmark":"OpenCompass — AIME2025","benchmark_slug":"oc-aime2025","score":89,"unit":"%","category":"math","source":"OpenCompass","source_url":"https://rank.opencompass.org.cn","benchmark_url":"https://benchgecko.ai/benchmark/oc-aime2025"},{"benchmark":"OpenCompass — MMLU-Pro","benchmark_slug":"oc-mmlu-pro","score":82,"unit":"%","category":"knowledge","source":"OpenCompass","source_url":"https://rank.opencompass.org.cn","benchmark_url":"https://benchgecko.ai/benchmark/oc-mmlu-pro"},{"benchmark":"HELM — IFEval","benchmark_slug":"helm-ifeval","score":81,"unit":"%","category":"language","source":"Stanford HELM","source_url":"https://crfm.stanford.edu/helm/","benchmark_url":"https://benchgecko.ai/benchmark/helm-ifeval"},{"benchmark":"HELM — WildBench","benchmark_slug":"helm-wildbench","score":80.7,"unit":"%","category":"reasoning","source":"Stanford HELM","source_url":"https://crfm.stanford.edu/helm/","benchmark_url":"https://benchgecko.ai/benchmark/helm-wildbench"},{"benchmark":"HELM — MMLU-Pro","benchmark_slug":"helm-mmlu-pro","score":78.6,"unit":"%","category":"knowledge","source":"Stanford HELM","source_url":"https://crfm.stanford.edu/helm/","benchmark_url":"https://benchgecko.ai/benchmark/helm-mmlu-pro"},{"benchmark":"OpenCompass — GPQA-Diamond","benchmark_slug":"oc-gpqa-diamond","score":77,"unit":"%","category":"knowledge","source":"OpenCompass","source_url":"https://rank.opencompass.org.cn","benchmark_url":"https://benchgecko.ai/benchmark/oc-gpqa-diamond"},{"benchmark":"LiveBench — Mathematics","benchmark_slug":"livebench-mathematics","score":74.26,"unit":"%","category":"math","source":"LiveBench","source_url":"https://livebench.ai","benchmark_url":"https://benchgecko.ai/benchmark/livebench-mathematics"},{"benchmark":"OpenCompass — LiveCodeBenchV6","benchmark_slug":"oc-livecodebenchv6","score":66.3,"unit":"%","category":"coding","source":"OpenCompass","source_url":"https://rank.opencompass.org.cn","benchmark_url":"https://benchgecko.ai/benchmark/oc-livecodebenchv6"},{"benchmark":"HELM — GPQA","benchmark_slug":"helm-gpqa","score":63,"unit":"%","category":"knowledge","source":"Stanford HELM","source_url":"https://crfm.stanford.edu/helm/","benchmark_url":"https://benchgecko.ai/benchmark/helm-gpqa"},{"benchmark":"LiveBench — Coding","benchmark_slug":"livebench-coding","score":60.66,"unit":"%","category":"coding","source":"LiveBench","source_url":"https://livebench.ai","benchmark_url":"https://benchgecko.ai/benchmark/livebench-coding"},{"benchmark":"LiveBench — Reasoning","benchmark_slug":"livebench-reasoning","score":58.16,"unit":"%","category":"reasoning","source":"LiveBench","source_url":"https://livebench.ai","benchmark_url":"https://benchgecko.ai/benchmark/livebench-reasoning"},{"benchmark":"LiveBench — Language","benchmark_slug":"livebench-language","score":56.31,"unit":"%","category":"language","source":"LiveBench","source_url":"https://livebench.ai","benchmark_url":"https://benchgecko.ai/benchmark/livebench-language"},{"benchmark":"LiveBench — Data Analysis","benchmark_slug":"livebench-data-analysis","score":53.58,"unit":"%","category":"reasoning","source":"LiveBench","source_url":"https://livebench.ai","benchmark_url":"https://benchgecko.ai/benchmark/livebench-data-analysis"},{"benchmark":"LiveBench — Overall","benchmark_slug":"livebench-overall","score":50.41,"unit":"%","category":"knowledge","source":"LiveBench","source_url":"https://livebench.ai","benchmark_url":"https://benchgecko.ai/benchmark/livebench-overall"},{"benchmark":"HELM — Omni-MATH","benchmark_slug":"helm-omni-math","score":46.7,"unit":"%","category":"math","source":"Stanford HELM","source_url":"https://crfm.stanford.edu/helm/","benchmark_url":"https://benchgecko.ai/benchmark/helm-omni-math"},{"benchmark":"LiveBench — If","benchmark_slug":"livebench-if","score":41.54,"unit":"%","category":"language","source":"LiveBench","source_url":"https://livebench.ai","benchmark_url":"https://benchgecko.ai/benchmark/livebench-if"},{"benchmark":"OpenCompass — HLE","benchmark_slug":"oc-hle","score":13.5,"unit":"%","category":"knowledge","source":"OpenCompass","source_url":"https://rank.opencompass.org.cn","benchmark_url":"https://benchgecko.ai/benchmark/oc-hle"},{"benchmark":"LiveBench — Agentic Coding","benchmark_slug":"livebench-agentic-coding","score":8.33,"unit":"%","category":"coding","source":"LiveBench","source_url":"https://livebench.ai","benchmark_url":"https://benchgecko.ai/benchmark/livebench-agentic-coding"}],"gecko_tests":null,"links":{"page":"https://benchgecko.ai/model/qwen3-next-80b-a3b-thinking","json":"https://benchgecko.ai/api/v1/models/qwen3-next-80b-a3b-thinking","price_history":"https://benchgecko.ai/api/v1/price-history/qwen3-next-80b-a3b-thinking","provider_prices":"https://benchgecko.ai/pricing/arbitrage/qwen3-next-80b-a3b-thinking"}}