Compare commits

...

306 Commits

Author SHA1 Message Date
Adam e1d75f89bc Normalize base-model inheritance 2026-06-04 12:52:03 -05:00
Adam c8a52d19c1 feat(models): add coding benchmarks and weights (#2000)
* feat(models): add coding benchmarks and weights

* feat(models): add more coding benchmarks

* feat(models): add agent benchmark scores

* feat(models): normalize benchmark metadata
2026-06-04 11:09:53 -05:00
Adam 859bb31ffa refactor(chutes): use base_model for tee wrappers (#2002) 2026-06-04 11:05:45 -05:00
Adam a672fe4b08 feat(web): surface model metadata links (#2001) 2026-06-04 11:03:53 -05:00
Frank 72fc81a4da update zen models 2026-06-04 11:31:06 -04:00
Aiden Cline 9e8cad9ac9 Merge pull request #1988 from anomalyco/feat/stepfun-reasoning-options
feat(stepfun): add reasoning effort options
2026-06-04 00:38:32 -05:00
Aiden Cline f6c1036a86 Merge pull request #1985 from anomalyco/feat/xiaomi-reasoning-options
feat(xiaomi): add reasoning toggles
2026-06-04 00:37:58 -05:00
Aiden Cline cb22ad5625 Merge pull request #1992 from anomalyco/fix/sync-base-model-output
fix(sync): preserve base model output
2026-06-04 00:02:46 -05:00
Aiden Cline 909db75087 fix(sync): preserve base model output 2026-06-04 00:00:46 -05:00
Aiden Cline b551552f14 Merge pull request #1983 from anomalyco/feat/cohere-reasoning-options
feat(cohere): add reasoning options
2026-06-03 23:34:36 -05:00
Aiden Cline 0919062b40 Merge pull request #1989 from anomalyco/feat/google-gemini-reasoning-options
feat(google): add Gemini reasoning options
2026-06-03 23:05:44 -05:00
Aiden Cline 8662c63313 fix(google): retain deprecated Gemini reasoning metadata 2026-06-03 23:03:08 -05:00
Aiden Cline 36b808691a Merge pull request #1991 from shzdehmd/dev
chore(fireworks): update qwen3p6-plus limits to 262K context / 65K output
2026-06-03 22:21:52 -05:00
Ahmad Shahzad c8feebb7cd chore(fireworks): update qwen3p6-plus limits to 262K context / 65K output 2026-06-04 07:05:29 +05:00
Aiden Cline 259801fba5 fix(google): exclude unavailable Gemini 3 Pro preview 2026-06-03 18:28:25 -05:00
Aiden Cline ce35e18561 feat(stepfun): add reasoning effort options 2026-06-03 17:49:21 -05:00
Aiden Cline 133a0b0126 feat(google): add Gemini reasoning options 2026-06-03 17:49:09 -05:00
Aiden Cline 99406ae7df feat(xiaomi): add reasoning toggles 2026-06-03 17:48:22 -05:00
Aiden Cline 4f25170be2 feat(cohere): add reasoning options 2026-06-03 17:48:05 -05:00
Aiden Cline 6ae56b00a8 Merge pull request #1981 from anomalyco/feat/glm-coding-plan-reasoning-toggle
feat(glm): add Zhipu and coding plan reasoning toggles
2026-06-03 17:31:39 -05:00
Aiden Cline 2cb0d28e17 feat(zhipuai): add reasoning toggles 2026-06-03 17:27:57 -05:00
Aiden Cline 03d90aeacc Merge pull request #1982 from anomalyco/fix/sync-model-catalog-matrix
fix(sync): restore model catalog workflow
2026-06-03 17:26:22 -05:00
Aiden Cline eb7dbead75 fix(sync): restore model catalog workflow 2026-06-03 17:20:01 -05:00
Aiden Cline 36cbbfc577 feat(glm): add coding plan reasoning toggles 2026-06-03 16:56:42 -05:00
Aiden Cline d84e883ede Merge pull request #1980 from anomalyco/feat/deepseek-reasoning-options
feat(deepseek): add reasoning options
2026-06-03 16:09:22 -05:00
Aiden Cline dd9d6ff54e feat(deepseek): add reasoning options 2026-06-03 16:08:09 -05:00
Aiden Cline d7e19c7627 Merge pull request #1979 from anomalyco/feat/moonshot-reasoning-toggle
feat(moonshot): add reasoning toggle options
2026-06-03 15:46:39 -05:00
Aiden Cline ca3eb39fbd feat(moonshot): add reasoning toggle options 2026-06-03 15:45:15 -05:00
Aiden Cline d136e7b036 Merge pull request #1955 from eliasaronson/chore/mark-deprecated-models
chore: mark retired models as deprecated
2026-06-03 15:25:27 -05:00
Adam f6c6f04367 feat(models): add model metadata (#1974)
* feat(models): add model metadata

* feat(models): rename model metadata namespaces
2026-06-03 15:13:54 -05:00
Aiden Cline 1e9e4bbdab Merge pull request #1977 from Ardakilic/feat/nano-gpt-20260603
chore: sync nano-gpt models: 20260603
2026-06-03 15:02:42 -05:00
Arda Kılıçdağı efb553d287 chore: sync nano-gpt models: 20260603 2026-06-03 21:31:28 +03:00
Jack d9b83ca9cd Merge pull request #1975 from anomalyco/update/opencode-go-qwen3.7-plus
feat(opencode-go): add Qwen3.7 Plus model
2026-06-04 01:30:11 +08:00
Aiden Cline 25592e361a Merge pull request #1976 from jerome-benoit/feat/sap-ai-core-gpt-5.5
feat(sap-ai-core): add GPT-5.5
2026-06-03 12:28:07 -05:00
Jérôme Benoit 012928800c fix(sap-ai-core): align GPT-5.4 release_date with upstream
SAP AI Core routes to OpenAI gpt-5.4; release_date should reflect
the actual model release (2026-03-05) rather than the SAP catalog
availability date (2026-04-27).
2026-06-03 19:09:54 +02:00
Jérôme Benoit ee60c2e4d4 feat(sap-ai-core): add GPT-5.5
SAP AI Core routes to OpenAI gpt-5.5; specs mirror the canonical
provider/openai/gpt-5.5 with the established sap-ai-core wrapper
adjustments (lowercase name, drop [[cost.tiers]], drop
[experimental.modes.fast]).
2026-06-03 19:06:07 +02:00
Jack 7e76cde0b1 fix(opencode-go): correct Qwen3.7 Plus dates 2026-06-04 00:59:42 +08:00
Jack bbd2479e12 feat(opencode-go): add Qwen3.7 Plus model 2026-06-04 00:57:42 +08:00
Aiden Cline eeb17ccab4 Merge pull request #1971 from coder-wangbin/feat/alibaba-cn-qwen3.7-plus
feat: add Qwen3.7 Plus model for alibaba-cn provider
2026-06-03 10:00:39 -05:00
Aiden Cline 46ffeee012 Merge pull request #1972 from Phosmachina/feature/update-deepinfra-deepseek-and-mimo
Feature/update deepinfra deepseek and mimo
2026-06-03 09:56:37 -05:00
Aiden Cline 8c1e45007d Merge pull request #1967 from anomalyco/automation/sync-models-openrouter
chore(sync): update OpenRouter model catalog
2026-06-03 09:54:14 -05:00
Frank 364628cb93 update zen models 2026-06-03 10:32:41 -04:00
github-actions[bot] 50f5d40b4e chore(sync): update OpenRouter model catalog 2026-06-03 14:02:08 +00:00
Michel Paronnaud 707be517c6 chore(deepinfra): adjust prices and limit for DeepSeek models 2026-06-03 12:40:49 +02:00
Michel Paronnaud 657598e609 fix(deepinfra): path for mimo models 2026-06-03 12:28:08 +02:00
wangbin 0999a475ba feat: add Qwen3.7 Plus model for alibaba-cn provider
- Add base definition in providers/alibaba/models/qwen3.7-plus.toml
- Add extends reference in providers/alibaba-cn/models/qwen3.7-plus.toml
- Release date: 2026-06-02
- Context: 131K tokens, Output: 16K tokens
- Pricing: $0.50/$3.00 per 1M tokens (input/output)
- Supports reasoning and tool calling
2026-06-03 17:40:12 +08:00
Frank e1ea0254ed update zen models 2026-06-02 22:51:19 -04:00
Aiden Cline 7fad18b054 fix 2026-06-02 16:36:29 -05:00
Aiden Cline 710bbc5375 Merge pull request #1961 from anomalyco/feat/minimax-m3-reasoning-toggle
feat(minimax): add M3 reasoning toggle
2026-06-02 16:34:19 -05:00
Aiden Cline 0b9fa62153 Merge pull request #1963 from nicholasgriffintn/open-mistral-nemo
fix: update mistral nemo
2026-06-02 14:34:54 -05:00
Aiden Cline 126a481a70 Merge pull request #1965 from nicholasgriffintn/update-devstral-models
chore: update devstral models
2026-06-02 14:34:38 -05:00
Aiden Cline b080cf1fe4 Merge pull request #1952 from anomalyco/automation/sync-models-xai
chore(sync): update xAI model catalog
2026-06-02 14:03:02 -05:00
Nicholas Griffin 042002d3a3 chore: update devstral models 2026-06-02 19:47:17 +01:00
Nicholas Griffin a192b51e84 chore: undo 2026-06-02 19:46:13 +01:00
Nicholas Griffin b870c796ef chore: undo 2026-06-02 19:45:25 +01:00
Nicholas Griffin 35e4981aa9 chore: update 2026-06-02 19:40:02 +01:00
Frank 970b660070 sync 2026-06-02 14:24:55 -04:00
Nicholas Griffin bb03a8c207 fix: update mistral nemo 2026-06-02 19:22:36 +01:00
github-actions[bot] b6dbb68c32 chore(sync): update xAI model catalog 2026-06-02 18:21:19 +00:00
Aiden Cline b34a7c21f3 feat(minimax): add M3 reasoning toggle 2026-06-02 12:43:17 -05:00
Aiden Cline 4b169b8cb6 Merge pull request #1960 from anomalyco/feat/zai-reasoning-toggle
feat(zai): add reasoning toggle options
2026-06-02 12:38:15 -05:00
Aiden Cline f8ecf4daec feat(zai): add reasoning toggle options 2026-06-02 12:02:54 -05:00
Frank eb68be3fbc stats 2026-06-02 12:38:49 -04:00
Aiden Cline 55fb058cc5 tweak: update available options 2026-06-02 11:35:26 -05:00
Aiden Cline 5545a98866 tweak: handle missing omits gracefully 2026-06-02 11:25:55 -05:00
Aiden Cline 1a7b103bba fix: omit 2026-06-02 11:22:52 -05:00
Aiden Cline f341858398 Merge pull request #1896 from stevenyeung/add/alibaba-token-plan
feat: add Alibaba Token Plan provider with 15 models
2026-06-02 11:20:42 -05:00
Aiden Cline f2d1e42582 Merge pull request #1951 from anomalyco/automation/sync-models-openrouter
chore(sync): update OpenRouter model catalog
2026-06-02 11:03:27 -05:00
Aiden Cline 2f9d2d8d78 Merge pull request #1957 from kapelame/fix/minimax-m3-prune
fix(minimax): correct MiniMax-M3 pricing and max output
2026-06-02 10:53:58 -05:00
Aiden Cline bf6b0e32d7 Merge pull request #1959 from Tavernari/feat/claudinio-add-audio-video
feat(claudinio): add audio and video input modalities
2026-06-02 10:53:25 -05:00
github-actions[bot] 5b4d2ea2e4 chore(sync): update OpenRouter model catalog 2026-06-02 15:41:53 +00:00
Victor Carvalho Tavernari ee4badab7b fix(claudinio): update cache_read price to 0.150 per 1M tokens 2026-06-02 16:34:05 +01:00
Victor Carvalho Tavernari 6c05a6b519 Merge branch 'dev' into feat/claudinio-add-audio-video 2026-06-02 16:32:12 +01:00
Victor Carvalho Tavernari 91ae9a48e3 feat(claudinio): add audio and video input modalities 2026-06-02 16:24:06 +01:00
kapelame 85d02711bd fix(minimax): correct MiniMax-M3 pricing and max output
The M3 entries added in #1940 copied M2.7's cost values. Correct them to
the official M3 pricing and limits:

- minimax / minimax-cn (pay-as-you-go): input 0.30 -> 0.60,
  output 1.20 -> 2.40, cache_read 0.06 -> 0.12, and remove cache_write
  (M3 has no active prompt-cache-write tier).
- max output 131072 -> 128000 across all four providers.
- coding-plan variants keep their subscription-plan zero pricing; only
  max output is corrected.

Context (512K), modalities, and the other flags are unchanged.
2026-06-02 21:03:45 +08:00
Elias H Aronsson 1ba404612d chore: mark retired models as deprecated
Add status = "deprecated" to models that are past their provider's
shutdown/retirement date (no longer served by the public API).

Google Gemini (4):
  gemini-2.0-flash, gemini-2.0-flash-lite, gemini-3-pro-preview,
  gemini-3.1-flash-lite-preview

Anthropic Claude (7):
  claude-3-sonnet-20240229, claude-3-5-sonnet-20240620,
  claude-3-5-sonnet-20241022, claude-3-opus-20240229,
  claude-3-7-sonnet-20250219, claude-3-5-haiku-20241022,
  claude-3-haiku-20240307

OpenAI (2):
  o1-preview, o1-mini

Sources:
  https://ai.google.dev/gemini-api/docs/deprecations
  https://platform.claude.com/docs/en/about-claude/model-deprecations
  https://developers.openai.com/api/docs/deprecations
2026-06-02 10:06:18 +02:00
Aiden Cline f91dd4ad0b Merge pull request #1950 from anomalyco/fix/reasoning-options-inheritance
fix: do not inherit reasoning options
2026-06-01 23:24:13 -05:00
Aiden Cline 449b926f40 fix: do not inherit reasoning options 2026-06-01 23:22:04 -05:00
Aiden Cline 2546ffe570 Merge pull request #1940 from matstrange/add-minimax-m3
Add MiniMax-M3 model (#1933)
2026-06-01 22:57:45 -05:00
Hex Agent f33ff9ba78 Add MiniMax-M3 model to 5 providers
MiniMax-M3 is MiniMax's new frontier multimodal coding model: 1M context
window (512K minimum on ollama-cloud), native text/image/video input,
tool calling, reasoning, and open weights.

Adds the model to all five providers where it should be available:

  - minimax (pay-as-you-go)
  - minimax-cn (pay-as-you-go, China)
  - minimax-coding-plan (token plan subscription)
  - minimax-cn-coding-plan (token plan subscription, China)
  - ollama-cloud

Closes #1933.

Notes for reviewers:
  - Cost fields on minimax/minimax-cn match M2.7; M3 docs state the
    pricing is unchanged from M2.7.
  - The 1M/512K context divergence on ollama-cloud is intentional —
    ollama advertises 1M with a 512K minimum, so 512K is the safe floor
    that won't surprise opencode users with mid-request rejections.
  - output = 131072 is inherited from the existing M2.7 files; MiniMax's
    published M3 docs only advertise the 1M input context, not a separate
    output cap.
  - The minimax-coding-plan variant has been verified end-to-end in
    opencode against the MiniMax token plan API.
2026-06-01 22:54:17 -05:00
Aiden Cline 98bf80cd77 Merge pull request #1949 from anomalyco/fix/github-copilot-context-limits-complete
fix(github-copilot): preserve model-specific limits
2026-06-01 22:49:00 -05:00
Aiden Cline 68660cad83 Merge pull request #1938 from anyapi-ai/dev
Add AnyAPI provider
2026-06-01 22:47:26 -05:00
Aiden Cline 78e6c1a2dd fix(github-copilot): preserve model-specific limits 2026-06-01 22:46:56 -05:00
Aiden Cline 062ba1fbb7 Merge pull request #1939 from anomalyco/automation/sync-models-openrouter
chore(sync): update OpenRouter model catalog
2026-06-01 22:46:18 -05:00
Aiden Cline 477ccffee8 Merge pull request #1948 from anomalyco/feat/anthropic-reasoning-options
feat(anthropic): add reasoning options
2026-06-01 22:41:23 -05:00
github-actions[bot] 6ce2f88eac chore(sync): update OpenRouter model catalog 2026-06-02 03:26:49 +00:00
Aiden Cline 4501589a20 feat(anthropic): add reasoning options 2026-06-01 22:17:07 -05:00
Aiden Cline 0b643efd17 Merge pull request #1947 from LYY/fix/github-copilot-context-limits
fix(github-copilot): correct limits for 8 models with verified CAPI data (partial, refs #1946)
2026-06-01 22:16:07 -05:00
LYY 8a0ea15aca fix(github-copilot): correct limits for models with verified CAPI data
The github-copilot model stubs use [extends] and inherit the upstream
base model's context/input/output limits (~1M), which do not match the
limits GitHub Copilot actually enforces via api.githubcopilot.com/models.

Override the limit fields for the 8 models whose live CAPI values were
verified first-hand. Models that were not enabled on the test account
(no CAPI data available) are intentionally left unchanged.

Refs anomalyco/models.dev#1946
2026-06-02 11:08:51 +08:00
Aiden Cline b3676fd77d Merge pull request #1934 from yukoba/github-copilot-2026-06
June 2026 changes of GitHub Copilot
2026-06-01 16:10:45 -05:00
Aiden Cline f97a02078c Merge pull request #1932 from eliasto/ovhcloud/update-models-qwen
feat(ovhcloud): Add new Qwen models and sync mode
2026-06-01 14:27:58 -05:00
Aiden Cline 2a7153fcef Merge pull request #1943 from peculiarnewbie/fix/crof-models-update
fix: update crof.ai models to match current API
2026-06-01 13:41:23 -05:00
Aiden Cline 3396f15854 Merge pull request #1942 from anomalyco/feat/reasoning-options-schema
feat: add reasoning options schema
2026-06-01 13:41:10 -05:00
Aiden Cline ccfd99b5e9 Merge pull request #1941 from KTibow/fix-llmgateway-deepseek-v3-2-price
fix(llmgateway): update model pricing
2026-06-01 13:40:20 -05:00
bolt ef6eb88216 fix: update crof.ai models to match current API 2026-06-02 00:52:14 +07:00
Aiden Cline 94e128244a feat: add reasoning options schema 2026-06-01 12:48:11 -05:00
KTibow ec6d07d41c fix(llmgateway): use weighted pricing selection 2026-06-01 10:41:23 -07:00
KTibow 4efbe58af8 fix(llmgateway): update model pricing 2026-06-01 10:24:05 -07:00
Christina 3b55102a45 Add AnyAPI provider with 30 models
Adds AnyAPI (https://anyapi.ai) as a new provider. Models reuse existing
canonical entries through `extends`. Cost fields are omitted as AnyAPI uses
a credit-based pricing system. Validated locally with `bun validate`.

Models (30):
- openai: gpt-5.4, gpt-5.2, gpt-5.1, gpt-5, gpt-5-mini, gpt-4.1, gpt-4.1-mini, o4-mini, o3, o3-mini
- anthropic: claude-opus-4-7, claude-opus-4-6, claude-sonnet-4-6, claude-sonnet-4-5, claude-haiku-4-5
- google: gemini-2.5-pro, gemini-2.5-flash, gemini-2.5-flash-lite, gemini-3-pro-preview, gemini-3-flash-preview
- deepseek: deepseek-v4-pro, deepseek-v4-flash, deepseek-chat, deepseek-r1
- mistralai: mistral-large-2512, devstral-2512
- perplexity: sonar-pro, sonar-reasoning-pro
- cohere: command-r-plus-08-2024
- xai: grok-4.3
2026-06-01 17:14:24 +02:00
Aiden Cline a4fe8fc36b fmt 2026-06-01 09:53:16 -05:00
Aiden Cline b24b44ca0a Merge pull request #1920 from anomalyco/automation/sync-models-openrouter
chore(sync): update OpenRouter model catalog
2026-06-01 09:39:43 -05:00
Aiden Cline 466d664bf9 Merge pull request #1930 from dpuyosa/dev
Venice: Add MiniMax M3 model
2026-06-01 09:39:27 -05:00
Aiden Cline 340ef131e9 Merge pull request #1936 from JoshuaDietz/dev
feat(ollama-cloud): add minimax m3
2026-06-01 09:38:47 -05:00
Aiden Cline 6ce69a706a Merge pull request #1937 from KTibow/fix/hpc-ai-pricing
fix: update HPC-AI model pricing
2026-06-01 09:38:34 -05:00
KTibow 3505674a1b fix: update HPC-AI model pricing 2026-06-01 07:29:43 -07:00
github-actions[bot] cfeaef8f77 chore(sync): update OpenRouter model catalog 2026-06-01 14:25:41 +00:00
Joshua Dietz f73a5d571b feat(ollama-cloud): add minimax m3 2026-06-01 16:14:49 +02:00
Yu Kobayashi 62d31abff8 June 2026 changes of GitHub Copilot 2026-06-01 22:03:44 +09:00
Elias TOURNEUX ee71ea7181 feat(ovhcloud): Add qwen 3.6 27b 2026-06-01 11:57:58 +02:00
Elias TOURNEUX b038f40b64 feat(ovhcloud): Add OVHcloud sync mode 2026-06-01 10:28:46 +02:00
Elias TOURNEUX 4b7bdd9692 feat(ovhcloud): Add new Qwen models 2026-06-01 10:16:54 +02:00
dpuyosa b73ef1a634 [venice] Add MiniMax M3 model
- Add minimax-m3.toml with cost, limits, modalities
- Enable text/image input and text output support
- Set 500k context, 32k output, cache pricing
2026-06-01 09:38:12 +02:00
Aiden Cline c11b5bb00b Merge pull request #1927 from Jercik/codex/update-wafer-price-cuts
fix: update Wafer.ai pricing
2026-06-01 00:01:03 -05:00
Łukasz Jerciński 01bc1785b4 fix: update Wafer price cuts 2026-06-01 06:25:20 +02:00
Aiden Cline 36b5847fa4 Merge pull request #1925 from anomalyco/fix/vercel-claude-opus-4-1-symlink
fix vercel claude opus 4.1 symlink
2026-05-31 22:12:27 -05:00
Aiden Cline 040beaf818 fix vercel claude opus 4.1 symlink 2026-05-31 22:11:05 -05:00
Frank 9f8e1a3858 update zen models 2026-05-31 22:07:52 -04:00
Jack ce7a1c0218 update context limit of M3 in Go 2026-06-01 09:45:46 +08:00
Jack 55e6c7a2ee add image,video modalities to M3 2026-06-01 08:58:02 +08:00
Frank 8bacf496fe update zen models 2026-05-31 19:49:41 -04:00
Frank ac12cf1f43 update zen models 2026-05-31 19:47:33 -04:00
Frank 06c7ea2303 update zen models 2026-05-31 19:47:04 -04:00
Frank 92b6e9da53 update zen models 2026-05-31 19:43:09 -04:00
Frank 660ff026a0 update zen models 2026-05-31 19:40:57 -04:00
Aiden Cline 9a910a1fb3 Merge pull request #1917 from markgibaud/fix/eu-opus-bedrock-pricing
fix(bedrock): correct EU cross-region pricing for Claude Opus and Sonnet models
2026-05-31 17:37:45 -05:00
Aiden Cline 7dc5f787b0 Merge pull request #1915 from anomalyco/automation/sync-models-openrouter
chore(sync): update OpenRouter model catalog
2026-05-31 17:10:10 -05:00
Aiden Cline c25ddce07c Merge pull request #1916 from anomalyco/automation/sync-models-cloudflare-workers-ai
chore(sync): update Cloudflare Workers AI model catalog
2026-05-31 17:05:57 -05:00
github-actions[bot] e73a82e44e chore(sync): update OpenRouter model catalog 2026-05-31 21:36:49 +00:00
github-actions[bot] 7de9710b5a chore(sync): update Cloudflare Workers AI model catalog 2026-05-31 21:36:49 +00:00
Frank d1a8dbcdde update zen models 2026-05-31 14:19:49 -04:00
markgibaud 556ef84d11 fix(bedrock): correct EU cross-region pricing for Claude Opus and Sonnet models
AWS Bedrock EU (Europe/London) cross-region inference has a 10% premium
over the base Anthropic pricing. The EU models were incorrectly using
the same pricing as the US/base models.

Affected models:
- eu.anthropic.claude-opus-4-6-v1
- eu.anthropic.claude-opus-4-7
- eu.anthropic.claude-opus-4-8
- eu.anthropic.claude-sonnet-4-5-20250929-v1:0
- eu.anthropic.claude-sonnet-4-6

Corrected Opus per 1M token prices:
- Input: $5.00 -> $5.50
- Output: $25.00 -> $27.50
- Cache read: $0.50 -> $0.55
- Cache write: $6.25 -> $6.875

Corrected Sonnet per 1M token prices:
- Input: $3.00 -> $3.30
- Output: $15.00 -> $16.50
- Cache read: $0.30 -> $0.33
- Cache write: $3.75 -> $4.125

Source: AWS Bedrock pricing page, Europe (London) region
2026-05-31 11:07:18 +01:00
Aiden Cline f2020553ea Merge pull request #1898 from Suat-B/codex/xpersona-frieren-1
Rename Xpersona Frieren display name
2026-05-30 16:44:39 -05:00
Aiden Cline ca0a7e17cb Merge pull request #1910 from anomalyco/automation/sync-models-openrouter
chore(sync): update OpenRouter model catalog
2026-05-30 16:44:22 -05:00
github-actions[bot] d2468eafb2 chore(sync): update OpenRouter model catalog 2026-05-30 21:37:16 +00:00
Aiden Cline 96d0f9beff Merge pull request #1913 from orangeclk/add-glm-4.6v-zhipuai-coding-plan
feat(zhipuai-coding-plan): add glm-4.6v symlink from zai
2026-05-30 11:52:50 -05:00
Aiden Cline 7989be9f34 Merge pull request #1912 from Jercik/codex/update-wafer-provider-models
feat: update Wafer provider models
2026-05-30 11:52:41 -05:00
Łukasz Jerciński f30e44f4b5 feat: update Wafer provider models 2026-05-30 14:53:02 +02:00
OrangeCLK 1383177893 feat(zhipuai-coding-plan): add glm-4.6v symlink from zai 2026-05-30 18:42:09 +08:00
Aiden Cline 2e58165af9 Merge pull request #1902 from kameshsampath/feat/provider/snowflake-cortex
feat(snowflake-cortex): add Snowflake Cortex provider
2026-05-29 23:55:55 -05:00
Kamesh Sampath f3466affc0 feat(snowflake-cortex): add Snowflake Cortex provider
Adds the snowflake-cortex provider which exposes Snowflake's Cortex
REST API (OpenAI Chat Completions-compatible endpoint) to opencode.

Provider details:
- npm: @ai-sdk/openai-compatible
- API: https://${SNOWFLAKE_ACCOUNT}.snowflakecomputing.com/api/v2/cortex/v1
- Auth: SNOWFLAKE_ACCOUNT + SNOWFLAKE_CORTEX_PAT (Programmatic Access Token)

Models (11, all with tool_call support):
- Anthropic: claude-opus-4-7 (beta/preview), claude-sonnet-4-6,
  claude-sonnet-4-5, claude-haiku-4-5
- OpenAI: openai-gpt-5.4 (beta), openai-gpt-5.2, openai-gpt-5.1,
  openai-gpt-5 (beta), openai-gpt-5-mini (beta), openai-gpt-5-nano (beta),
  openai-gpt-4.1

Models without tool_call support (deepseek-r1, llama3.1-70b,
snowflake-llama-3.3-70b, mistral-large2) are excluded per Snowflake docs:
"Tool calling is supported for OpenAI and Claude models only."

Preview/not-GA models are marked with status = "beta".

Cost fields are intentionally omitted on all models. The Cortex REST API
is billed in USD (AI_INFERENCE service type) at rates defined in the
Snowflake Service Consumption Table.

Closes #1895
2026-05-30 09:04:43 +05:30
Aiden Cline 3697c99297 Merge pull request #1901 from anomalyco/automation/sync-models-openrouter
chore(sync): update OpenRouter model catalog
2026-05-29 20:50:33 -05:00
Aiden Cline 796dda2fd1 Merge pull request #1904 from sdnts/dev
feat(cloudflare-ai-gateway): add opus 4.8
2026-05-29 20:50:11 -05:00
Aiden Cline 1d363cbd7f Merge pull request #1908 from anomalyco/fix/unique-provider-names
Ensure provider names are unique
2026-05-29 20:49:54 -05:00
Aiden Cline 3758725979 Ensure provider names are unique
Rename the China StepFun provider to "StepFun (China)" so it no longer
collides with the international "StepFun" provider, following the existing
Alibaba / Alibaba (China) naming convention.

Add a validation check in generate() that throws when two providers share
a (case-insensitive) name, preventing future duplicates.

Fixes #1906
2026-05-29 20:48:09 -05:00
github-actions[bot] c039e82f12 chore(sync): update OpenRouter model catalog 2026-05-30 01:16:42 +00:00
Siddhant 1e90242cdb feat(cloudflare-ai-gateway): add opus 4.8 2026-05-29 13:55:09 -04:00
Aiden Cline 277ac8577e Merge pull request #1897 from mikeyp/update-digitalocean-models
Add deepseek-4-flash and claude-opus-4.8 for DigitalOcean
2026-05-29 10:20:26 -05:00
Aiden Cline 6e5f35546d Merge pull request #1900 from ceyhanmolla/add-step-3.7-flash-v2
Add Step 3.7 Flash model to NVIDIA provider
2026-05-29 10:20:02 -05:00
ceyhanmolla aa69da14d6 Add Step 3.7 Flash model to NVIDIA provider
StepFun AI's Step 3.7 Flash - sparse MoE multimodal reasoning model:
- 198B total params, ~11B active per token
- 256K context window with sliding window attention
- Text + image input, text output
- Reasoning, tool calling, and attachment support
- Apache 2.0 license
2026-05-29 12:55:13 +02:00
Jack c7af5ba9be update model in Go 2026-05-29 16:20:16 +08:00
Suat-B d6e9e6735f Rename Xpersona Frieren model display name 2026-05-29 03:00:21 -05:00
Mike Prasuhn a987719410 Add deepseek-4-flash and claude-opus-4.8 2026-05-29 03:45:50 -04:00
Steven Yeung 3b792029c3 feat(alibaba-token-plan): add provider and 15 model TOMLs
Add Alibaba Cloud Model Studio Token Plan (Team Edition) provider with:
- Provider config (Singapore region, OpenAI-compatible endpoint)
- 10 models using extends pattern (zero-cost overrides from canonical providers)
- 5 models with full definitions (no canonical source available)
- Image generation models (qwen-image, wan2.7) with output=0 per convention
- deepseek-v3.2 with structured_output and corrected release date

Models: qwen3.7-max, qwen3.6-flash, qwen3.6-plus, kimi-k2.5, kimi-k2.6,
glm-5, glm-5.1, MiniMax-M2.5, deepseek-v4-pro, deepseek-v4-flash,
deepseek-v3.2, qwen-image-2.0, qwen-image-2.0-pro, wan2.7-image, wan2.7-image-pro
2026-05-29 15:12:27 +08:00
Aiden Cline 9528528f69 Merge pull request #1890 from anomalyco/automation/sync-models-openrouter
chore(sync): update OpenRouter model catalog
2026-05-28 23:42:03 -05:00
github-actions[bot] fbad40d863 chore(sync): update OpenRouter model catalog 2026-05-29 03:26:05 +00:00
Aiden Cline 84559b598b Merge pull request #1894 from anomalyco/add-siliconflow-cn-deepseek-v4-pro
feat(siliconflow-cn): add deepseek-ai/DeepSeek-V4-Pro
2026-05-28 19:10:18 -05:00
Aiden Cline 9d45730db9 Merge pull request #1889 from fhennerkes/dev
poe: add Claude-Opus-4.8 model
2026-05-28 19:10:06 -05:00
Aiden Cline 791089e4aa Merge pull request #1891 from ticoombs/dev
feat(copilot): update all copilot models, add claude-opus-4.8
2026-05-28 19:09:52 -05:00
Aiden Cline fcc4387187 feat(siliconflow-cn): add deepseek-ai/DeepSeek-V4-Pro 2026-05-28 19:09:15 -05:00
Tim C 6aa32ba566 fix(copilot): update all context,input,output with correct limits 2026-05-29 08:56:19 +10:00
Tim C 459a563d2b feat(copilot): add claude-opus-4.8 2026-05-29 08:52:42 +10:00
Aiden Cline 739e5a7c8e Merge pull request #1877 from aakash-gupte/add-merge-gateway-provider
Add Merge Gateway provider
2026-05-28 16:55:40 -05:00
Aiden Cline efc87afa3a Merge pull request #1888 from smakosh/add-llmgateway-opus-4-8
feat(llmgateway): add Claude Opus 4.8
2026-05-28 16:54:46 -05:00
Aiden Cline f616aa6a0b Merge pull request #1887 from dpuyosa/dev
Venice: Add Claude Opus 4.8 models
2026-05-28 16:54:35 -05:00
fhennerkes f77e75a567 poe: add Claude-Opus-4.8 model
Add new Anthropic model from Poe API (released 2026-05-28).
Uses extends format inheriting from anthropic/claude-opus-4-8
with Poe-specific overrides (name format, markup pricing,
slightly different context limit).

Co-Authored-By: Claude Opus 4.7 <noreply@anthropic.com>
2026-05-28 14:39:52 -07:00
smakosh 651bd4d91a feat(llmgateway): add Claude Opus 4.8
Extends anthropic/claude-opus-4-8, omitting the fast mode which LLM
Gateway does not expose, matching the existing 4.6/4.7 entries.

Co-Authored-By: Claude Opus 4.8 (1M context) <noreply@anthropic.com>
2026-05-28 22:07:40 +01:00
dpuyosa 30b9170106 [venice] Add Claude Opus 4.8 models
- Add claude-opus-4-8.toml (standard pricing)
- Add claude-opus-4-8-fast.toml (fast variant)
- Both: 1M context, image input, Opus 4.8 family
2026-05-28 22:54:39 +02:00
Aiden Cline 81e42c653c Merge pull request #1886 from vglafirov/add-gitlab-opus-4-8
feat: add GitLab Agentic Chat Opus 4.8
2026-05-28 15:25:06 -05:00
Vladimir Glafirov 11f2a38594 feat: add GitLab Agentic Chat Opus 4.8 2026-05-28 22:12:12 +02:00
Aiden Cline 6621be978f Merge pull request #1880 from bas3line/update-routing-run-base-url
Update routing.run API base URL
2026-05-28 14:46:09 -05:00
Aiden Cline c8d72661d6 Merge pull request #1885 from alaviss/bedrock-opus-4-8
feat: add cross-region inference entries for Bedrock for Opus 4.8
2026-05-28 14:45:43 -05:00
Hiếu Lê 21b7cacc33 feat: add cross-region inference entries for Bedrock for Opus 4.8 2026-05-28 12:38:19 -07:00
Aiden Cline ac1dc14439 Merge pull request #1884 from calebboyd/update-vercel-models-latest
feat: add vercel ai gateway opus 4.8
2026-05-28 14:31:03 -05:00
Aiden Cline 63acb799da Merge pull request #1883 from anomalyco/automation/sync-models-openrouter
chore(sync): update OpenRouter model catalog
2026-05-28 14:30:56 -05:00
github-actions[bot] 2e6ebb6750 chore(sync): update OpenRouter model catalog 2026-05-28 19:11:47 +00:00
calebboyd d8c2e15632 feat: add vercel ai gateway opus 4.8 2026-05-28 13:33:37 -05:00
Frank 8a70160200 update zen model 2026-05-28 14:01:35 -04:00
Aiden Cline 7ea1d4e15c Merge pull request #1882 from anomalyco/add-opus-4.8
feat: add opus 4.8
2026-05-28 12:04:51 -05:00
Aiden Cline 1c71b1bf17 feat: add opus 4.8 2026-05-28 12:04:14 -05:00
Aiden Cline acc3126194 Merge pull request #1879 from Ardakilic/chore/ci-fork-ensurance
CI Scheduled fork upstream ensurance
2026-05-28 11:27:47 -05:00
Aiden Cline 1f67addc01 Merge pull request #1878 from anomalyco/automation/sync-models-openrouter
chore(sync): update OpenRouter model catalog
2026-05-28 11:27:07 -05:00
github-actions[bot] bb011e71a7 chore(sync): update OpenRouter model catalog 2026-05-28 15:28:44 +00:00
bas3line 3df181411f fix(routing-run): update API base URL 2026-05-28 08:58:33 +05:30
Arda Kilicdagi a1f588463c chore: prevent forks to run cronjobbed sync commands 2026-05-28 06:25:45 +04:00
Arda Kilicdagi cb414c6886 chore: prevent forks to run cronjobbed sync commands 2026-05-28 05:31:19 +04:00
Arda Kilicdagi 4095e819c9 chore: prevent forks to run cronjobbed sync commands
chore: prevent forks to run cronjobbed sync commands
2026-05-28 05:26:15 +04:00
Aakash Gupte 9367cc1825 Add Merge Gateway logo
Co-Authored-By: Claude Opus 4.7 (1M context) <noreply@anthropic.com>
2026-05-27 17:22:44 -04:00
Aakash Gupte 96b9a04c4b Add Merge Gateway provider
Merge Gateway (https://merge.dev) is an LLM gateway exposing an
OpenAI/Anthropic-compatible API across many providers, using the
published `merge-gateway-ai-sdk-provider` npm package. Models are
defined via `extends` from existing canonical entries.

Co-Authored-By: Claude Opus 4.7 (1M context) <noreply@anthropic.com>
2026-05-27 14:32:21 -04:00
Jack d9de217dc4 update zen models 2026-05-28 01:46:38 +08:00
Aiden Cline 64ea80d416 Merge pull request #1869 from anomalyco/fix/xiaomi-token-plan-models
fix Xiaomi Token Plan model catalog
2026-05-27 12:02:22 -05:00
Aiden Cline d3859517d1 Merge pull request #1870 from anomalyco/automation/sync-models-openrouter
chore(sync): update OpenRouter model catalog
2026-05-27 12:01:59 -05:00
Aiden Cline 985a144c5f Merge pull request #1866 from tarikko/patch-1
Update Mistral Small model details in TOML file
2026-05-27 11:43:52 -05:00
Jack 5b05070eb4 Merge pull request #1875 from anomalyco/update/opencode-go-mimo-v2-5-pricing
update opencode go MiMo V2.5 pricing
2026-05-28 00:39:53 +08:00
Aiden Cline dad245c5e1 Merge pull request #1871 from oskarkocol/chore/update-groq-cache-prices
chore: update groq cache pricing
2026-05-27 11:39:22 -05:00
Aiden Cline bcff9bd6b2 Merge pull request #1874 from Alex-yang00/codex/sync-novita-models
Sync NovitaAI models
2026-05-27 11:39:03 -05:00
Aiden Cline 561bdaa546 Merge pull request #1876 from sebastiand-cerebras/deprecate-cerebras-qwen235b-llama8b-20260527
Remove deprecated Cerebras Qwen 3 235B model
2026-05-27 11:38:26 -05:00
Jack f4f210986c fix opencode go MiMo pricing conflict resolution 2026-05-28 00:36:48 +08:00
Jack f4206f4eaa Merge branch 'dev' into update/opencode-go-mimo-v2-5-pricing 2026-05-28 00:34:02 +08:00
Seb Duerr 52e5b67bc4 Remove deprecated Cerebras Qwen 3 235B model 2026-05-27 09:13:50 -07:00
Jack d11b0e448c update opencode go MiMo V2.5 pricing 2026-05-27 23:54:58 +08:00
github-actions[bot] 591745690e chore(sync): update OpenRouter model catalog 2026-05-27 15:28:06 +00:00
Codex 5a4bc9b383 Add selected NovitaAI models 2026-05-27 21:30:38 +08:00
Tarik a14171bcd4 Update mistral-small.toml
the only difference is the price so I removed the other fields
2026-05-27 14:00:33 +01:00
oskar 0741c5a53c chore: update kimi cache rate 2026-05-27 13:09:20 +07:00
oskar 7a1ec8737b chore: update groq cache pricing 2026-05-27 13:05:24 +07:00
Aiden Cline 7c37c92c2d fix Xiaomi Token Plan model catalog 2026-05-27 00:39:17 -05:00
Aiden Cline ec4ec6d441 Merge pull request #1867 from anomalyco/automation/sync-models-openrouter
chore(sync): update OpenRouter model catalog
2026-05-27 00:00:17 -05:00
Aiden Cline d96da5ece6 Merge pull request #1868 from Ardakilic/chore/sync-kilo-models-20260527-1
Sync Kilo models with upstream gateway
2026-05-26 23:59:26 -05:00
github-actions[bot] 0342c79d03 chore(sync): update OpenRouter model catalog 2026-05-27 03:26:29 +00:00
Arda Kilicdagi 37136fc2f3 chore: sync upstream kilo api gateway models 2026-05-27 02:27:53 +04:00
Tarik 131eca9c5e Update Mistral Small model configuration
use [extends] syntax
2026-05-26 20:44:22 +01:00
Tarik 3e38d46c4c Update Mistral Small model details in TOML file
The details on helicone website are false,

accurate details are pulled from deepinfra website

https://www.helicone.ai/model/mistral-small

https://deepinfra.com/mistralai/Mistral-Small-3.2-24B-Instruct-2506
2026-05-26 19:34:07 +01:00
Aiden Cline 989939773b Merge pull request #1863 from anomalyco/automation/sync-models-openrouter
chore(sync): update OpenRouter model catalog
2026-05-26 13:04:26 -05:00
Frank 4f1d5c511a update go models 2026-05-26 13:54:19 -04:00
github-actions[bot] 34a21e07f2 chore(sync): update OpenRouter model catalog 2026-05-26 17:26:26 +00:00
Aiden Cline 979f5da0ca Merge pull request #1864 from oskarkocol/update/openai-cache-rates
chore: update openai cache rates
2026-05-26 12:06:15 -05:00
oskar 73b357e2ab update cache rates 2026-05-26 23:28:11 +07:00
Aiden Cline 97c71f1663 Merge pull request #1862 from anomalyco/automation/sync-models-openrouter
chore(sync): update OpenRouter model catalog
2026-05-26 09:11:54 -05:00
github-actions[bot] 59c270bc85 chore(sync): update OpenRouter model catalog 2026-05-26 13:23:20 +00:00
Aiden Cline a4c18f88ed Merge pull request #1860 from bas3line/sync-routing-run-models-2
Sync routing.run model catalog
2026-05-25 23:44:24 -05:00
bas3line 3343a585ac fix(routing-run): sync model catalog 2026-05-26 09:51:14 +05:30
Aiden Cline f23db95550 Merge pull request #1858 from smakosh/fix/llmgateway-qwen3.7-max-id
fix(llmgateway): correct Qwen3.7 Max model id to qwen3.7-max
2026-05-25 17:22:56 -05:00
Aiden Cline 556d9a9045 Merge pull request #1859 from Suat-B/update-xpersona-frieren-coder-limits
Update Xpersona Frieren Coder limits
2026-05-25 17:22:46 -05:00
Frank 49991c8f8f update zen models 2026-05-25 17:56:52 -04:00
Suat-B 20c1e810ce Update Xpersona Frieren Coder limits 2026-05-25 13:23:59 -05:00
smakosh 15d08b4540 fix(llmgateway): correct Qwen3.7 Max model id to qwen3.7-max
Rename qwen37-max.toml to qwen3.7-max.toml so the model id matches
the canonical alibaba/qwen3.7-max definition (name: Qwen3.7 Max).

Co-Authored-By: Claude Opus 4.7 (1M context) <noreply@anthropic.com>
2026-05-25 18:48:12 +01:00
Aiden Cline 14bbc303e5 Merge pull request #1852 from anomalyco/automation/sync-models-openrouter
chore(sync): update OpenRouter model catalog
2026-05-25 09:49:26 -05:00
Aiden Cline 02da63ed48 Merge pull request #1853 from dpuyosa/chore/venice-pricing
Venice: Update pricing and limits for 4 models
2026-05-25 09:49:09 -05:00
github-actions[bot] 07daef8c54 chore(sync): update OpenRouter model catalog 2026-05-25 13:27:22 +00:00
dpuyosa b6ebe8696c [venice] Update pricing, limits, and last_updated for 4 models
- Reduce input/output/cache prices for gemini-3-5-flash, google-gemma-4-31b-it, and qwen-3-7-max
- Lower max output tokens from 65,536 to 16,384 for qwen3-5-35b-a3b
- Bump last_updated to 2026-05-25 for all 4 models
2026-05-25 10:07:32 +02:00
Aiden Cline ad654e71da Merge pull request #1850 from huxeon/dev
fix: modify the deepseek v4 flash/pro price
2026-05-25 00:18:18 -05:00
Aiden Cline 508f4d48e1 Merge pull request #1847 from ceyhanmolla/poolside/laguna-direct
Add Poolside provider with Laguna M.1 and XS.2 models
2026-05-24 23:26:34 -05:00
opencode-agent[bot] cd7c70b4fe revert: remove opencode-go deepseek-v4-pro price changes
Keep only the deepseek provider price updates as intended.
2026-05-25 04:26:14 +00:00
Aiden Cline 10e752ea84 Merge pull request #1845 from yukoba/vultr
Update Vultr models
2026-05-24 23:26:12 -05:00
huxeon 4cdb4c700b fix: modify the deepseek v4 flash/pro price in provider deepseek and opencode-go 2026-05-24 13:33:56 +08:00
Aiden Cline d497a446eb Merge pull request #1849 from technoabsurdist/add-wafer-ai-qwen3.6-35b-a3b-and-kimi-k2.6
providers/wafer.ai: add Qwen3.6-35B-A3B and Kimi-K2.6
2026-05-23 16:41:24 -05:00
Emilio Andere 745cf557b3 feat(wafer.ai): add Qwen3.6-35B-A3B and Kimi-K2.6
Both models are public serverless on pass.wafer.ai/v1/models but were
missing from the wafer.ai provider in models.dev, so OpenCode and other
tools that pull from the registry could not discover them.

- Qwen3.6 35B A3B: compact MoE, 32K context, vision-capable, $0.19/M in,
  $1.25/M out (NVFP4 on AMD MI355X — see wafer.ai/blog/qwen36-mi355x).
- Kimi K2.6: 1T sparse MoE, 262K context, vision-capable, $1.10/M in,
  $4.80/M out (NVFP4 on Blackwell — see wafer.ai/blog/kimi-k26-nvfp4).

Co-authored-by: Cursor <cursoragent@cursor.com>
2026-05-23 13:26:18 -04:00
Yu Kobayashi db295d6842 Refactor to use the extends syntax 2026-05-24 01:20:31 +09:00
ceyhanmolla 47c0eed41e Add Poolside provider with Laguna M.1 and XS.2 models 2026-05-23 17:09:33 +02:00
Yu Kobayashi e690980857 Update Vultr models 2026-05-23 16:25:11 +09:00
Aiden Cline f5f7d1a167 Merge pull request #1840 from anomalyco/automation/sync-models-openrouter
chore(sync): update OpenRouter model catalog
2026-05-22 16:58:23 -05:00
Aiden Cline cfbbb55c4d Merge pull request #1839 from anthraxx/alibaba-qwen3.7-plus
Alibaba: Add Qwen 3.7 Max and 3.6 Flash to all regions and plans
2026-05-22 16:58:08 -05:00
github-actions[bot] 1956762cdf chore(sync): update OpenRouter model catalog 2026-05-22 21:41:35 +00:00
Levente Polyak 28571433f7 feat(alibaba): add Qwen3.7 Max model configuration to all regions
Link: https://bailian.console.alibabacloud.com/cn-beijing?tab=model#/model-market/detail/qwen3.7-max?serviceSite=asia-pacific-china
2026-05-22 19:00:29 +02:00
Levente Polyak aef5e48bae feat(alibaba): add Qwen3.6 Flash model configuration to all regions
Link: https://bailian.console.alibabacloud.com/cn-beijing?tab=model#/model-market/detail/qwen3.6-flash?serviceSite=asia-pacific-china
2026-05-22 18:54:22 +02:00
Aiden Cline 8ba19639a7 Merge pull request #1825 from shzdehmd/dev
update(fireworks): sync models and pricing with current offerings
2026-05-22 11:45:45 -05:00
Aiden Cline 0f9b4c9edc Merge pull request #1834 from monotykamary/chore/update-neuralwatt-qwen3.6-pricing
fix(neuralwatt): update Qwen3.6 pricing to match API
2026-05-22 09:13:02 -05:00
Aiden Cline 2a9b6256dc Merge pull request #1836 from PierreLeGuen/nearai-provider
Add current NEAR AI Cloud models
2026-05-22 09:12:36 -05:00
Aiden Cline 51633fe106 Merge pull request #1838 from Quentinchampenois/fix/update-models-scaleway
fix: update Scaleway provider models list
2026-05-22 09:12:26 -05:00
Aiden Cline 169b5c4331 Merge pull request #1833 from NicoAvanzDev/add-copilot-gemini-3-5-flash
[GitHub Copilot] add Gemini 3.5 Flash
2026-05-22 09:09:25 -05:00
Aiden Cline 9d0f0d6d56 Merge pull request #1837 from fydrah/fix/google-vertex-gemini-3.5-flash
fix: missing google-vertex gemini 3.5 flash extend
2026-05-22 09:07:34 -05:00
Aiden Cline 1d2af0c97b Merge pull request #1835 from dpuyosa/feat/venice-models
Venice: Add Gemini 3.5 Flash and Qwen 3.7 Max, update Grok Build 0.1 pricing
2026-05-22 09:07:03 -05:00
Aiden Cline cc14ae4370 Merge pull request #1832 from anomalyco/automation/sync-models-openrouter
chore(sync): update OpenRouter model catalog
2026-05-22 09:06:16 -05:00
github-actions[bot] ad9a3448d2 chore(sync): update OpenRouter model catalog 2026-05-22 14:03:32 +00:00
Quentin Champenois c6f9cf58e4 fix(scaleway): add new models gemma-4-26b-a4b-it and qwen3.6-35b-a3b 2026-05-22 15:43:11 +02:00
Quentin Champenois bba5809471 fix(scaleway): Extends existing models 2026-05-22 15:42:36 +02:00
Quentin Champenois c0852f1b4f fix(scaleway): Clear removed models from list 2026-05-22 15:11:44 +02:00
Flavien Hardy 0fc3a3f635 fix: missing google-vertex gemini 3.5 flash extend 2026-05-22 08:56:44 -04:00
Pierre LE GUEN 68936dc062 Add current NEAR AI Cloud models 2026-05-22 09:27:57 +00:00
dpuyosa ade9760060 [venice] Add Gemini 3.5 Flash and Qwen 3.7 Max, update Grok Build 0.1 pricing
- Add Gemini 3.5 Flash model (1M context, multimodal input)
- Add Qwen 3.7 Max model (1M context, text-only)
- Update Grok Build 0.1 cost tiers and pricing
2026-05-22 10:57:38 +02:00
Tom X Nguyen 2a8b90b197 fix(neuralwatt): update Qwen3.6 pricing to match API
Updates the per-token cost for Qwen3.6-35B-A3B and qwen3.6-35b-fast from /bin/bash.05//bin/bash.10 to /bin/bash.29/.15 (input/output per million tokens), matching the actual Neuralwatt API pricing as reflected in pi-neuralwatt-provider commit f634286.
2026-05-22 15:19:51 +07:00
NicoAvanzDev 8569f0dfef [GitHub Copilot] add Gemini 3.5 Flash 2026-05-22 07:57:10 +00:00
Aiden Cline fc98ceb72e fix sync workflow force lease 2026-05-21 23:52:34 -05:00
Aiden Cline be7c5afc94 Merge pull request #1831 from anomalyco/fix-vertex-sonnet-4-6-limits
Fix Vertex Sonnet 4.6 token limits
2026-05-21 23:38:18 -05:00
Aiden Cline 57e62c43b0 fix vertex sonnet 4.6 limits 2026-05-21 23:37:30 -05:00
Aiden Cline 0898c35c9f Merge pull request #1830 from zainhas/dev
[Together AI] add Qwen3.7 max
2026-05-21 21:00:51 -05:00
Zain Hasan 46b23fb313 Merge branch 'anomalyco:dev' into dev 2026-05-21 17:47:35 -07:00
Zain Hasan 05fedc76cc [Together AI] add Qwen3.7 2026-05-21 17:47:19 -07:00
Aiden Cline 2738f81d1a Merge pull request #1828 from anomalyco/refactor/sync-core-layout
refactor: move sync implementation into core src
2026-05-21 18:10:09 -05:00
Aiden Cline 89b834086a refactor: move sync implementation into core src 2026-05-21 18:06:25 -05:00
Aiden Cline 1ab2ff8163 Merge pull request #1826 from smakosh/add-llmgateway-models
feat: add new LLM Gateway text models
2026-05-21 17:58:29 -05:00
Frank 9468676683 update zen models 2026-05-21 18:42:36 -04:00
Claude 6cdd2f054b Merge upstream/dev into add-llmgateway-models; resolve gemini-3.5-flash conflict
# Conflicts:
#	providers/google/models/gemini-3.5-flash.toml
2026-05-21 21:48:48 +00:00
Aiden Cline b13abc9141 Merge pull request #1827 from anomalyco/update-xai-pricing
fix xAI long-context pricing
2026-05-21 16:45:43 -05:00
Aiden Cline e5ba264751 fix xAI long-context pricing 2026-05-21 16:41:36 -05:00
smakosh a7811fb522 refactor: use extends for gemini and qwen models
Add canonical google/gemini-3.5-flash and alibaba/qwen3.7-max defs and
have the llmgateway entries extend them, per PR review.

Co-Authored-By: Claude Opus 4.7 (1M context) <noreply@anthropic.com>
2026-05-21 23:38:57 +02:00
smakosh 605fae75d9 feat: add new LLM Gateway text models
Add Grok 4.20 (reasoning/non-reasoning), Gemini 3.5 Flash, and Qwen3.7 Max to the llmgateway provider.

Co-Authored-By: Claude Opus 4.7 (1M context) <noreply@anthropic.com>
2026-05-21 23:04:47 +02:00
Ahmad Shahzad bcab0885bc update(fireworks): sync models and pricing with current offerings
Removed — 11 deprecated models:
- deepseek-v3p1
- deepseek-v3p2
- glm-4p5
- glm-4p5-air
- glm-4p7
- glm-5
- kimi-k2-instruct
- kimi-k2-thinking
- minimax-m2p1
- routers/kimi-k2p5-turbo

Added — 2 new Turbo (routers) models:
- routers/glm-5p1-fast
- routers/kimi-k2p6-turbo

Modified — pricing fixes:
- deepseek-v4-pro: cache_read 0.15 → 0.145
- gpt-oss-120b: added cache_read = 0.015
- gpt-oss-20b: input 0.05 → 0.07, output 0.20 → 0.30, added cache_read = 0.035
- minimax-m2p7: cache_read 0.03 → 0.06
2026-05-22 01:55:31 +05:00
Aiden Cline 26b05268ae Merge pull request #1824 from anomalyco/fix/vercel-gemini-35-flash
Add new Vercel AI Gateway models
2026-05-21 15:30:35 -05:00
Aiden Cline 1aee13d2e5 Add new Vercel AI Gateway models 2026-05-21 13:21:25 -05:00
Aiden Cline 0a924e6bb2 Merge pull request #1822 from anomalyco/automation/sync-models-xai
chore(sync): update xAI model catalog
2026-05-21 13:05:20 -05:00
Aiden Cline 9769b2b11d Merge pull request #1823 from anomalyco/automation/sync-models-openrouter
chore(sync): update OpenRouter model catalog
2026-05-21 13:05:13 -05:00
github-actions[bot] 1b0db099cf chore(sync): update OpenRouter model catalog 2026-05-21 17:56:13 +00:00
github-actions[bot] 16ed78587c chore(sync): update xAI model catalog 2026-05-21 17:56:11 +00:00
Frank 0b88965165 update zen models 2026-05-21 13:41:55 -04:00
Aiden Cline acc704ce39 Merge pull request #1797 from arnavchachra/add-crof-provider
add crof.ai provider with 21 models
2026-05-21 11:43:02 -05:00
Aiden Cline 51ad3b264e Merge pull request #1821 from anomalyco/sync-provider-ci
chore: automate provider sync jobs
2026-05-21 11:23:49 -05:00
Aiden Cline 146b6c7084 Merge pull request #1819 from Inceptron-Software/add_inceptron_provider
Add Inceptron provider
2026-05-21 11:21:27 -05:00
Aiden Cline 0e3cbe3c64 chore: automate provider sync jobs 2026-05-21 11:21:12 -05:00
Aiden Cline 604d4d66a4 Merge pull request #1820 from Suat-B/codex/xpersona-image-input-20260521
Add image input modality to Xpersona model
2026-05-21 11:00:49 -05:00
SuatB f5090028b8 Add image input modality to Xpersona model 2026-05-21 09:46:28 -05:00
Frank 4bad8faf29 update zen models 2026-05-21 09:05:13 -04:00
Oskar Gustafsson 0df2ccf586 Add Inceptron provider 2026-05-21 09:24:43 +02:00
Aiden Cline bafdc00b45 Merge pull request #1812 from anomalyco/openrouter-extends-sync
Sync OpenRouter models with extends
2026-05-20 21:07:21 -05:00
Aiden Cline 49840c013b Merge pull request #1814 from neonn0d/feat/stepfun-ai
feat(stepfun-ai): add international StepFun platform
2026-05-20 20:58:37 -05:00
Aiden Cline eccae0b54e sync openrouter models with extends 2026-05-20 20:32:11 -05:00
Aiden Cline 4cca29405f Merge pull request #1817 from dpuyosa/dev
Venice: Remove Grok 4.1 Fast and add Grok Build 0.1
2026-05-20 20:26:45 -05:00
Aiden Cline e40d9dd338 Merge pull request #1818 from anomalyco/cloudflare-sync-env
chore(sync): isolate cloudflare credentials
2026-05-20 20:26:22 -05:00
dpuyosa 035999cb58 [venice] Replace Grok 4.1 Fast with Grok Build 0.1
- Remove deprecated grok-41-fast model entry
- Add grok-build-0-1 with 200K token tiered pricing
- Update context to 256K and output limit to 65,536
2026-05-21 02:36:19 +02:00
Frank cec56bf1bc update zen models 2026-05-20 19:43:25 -04:00
Aiden Cline 85f0cdcb2f Merge pull request #1816 from anomalyco/xai-sync
Add PDF input modality to Grok models
2026-05-20 18:12:47 -05:00
Aiden Cline ef80d4df4e Infer PDF modality for xAI image models 2026-05-20 18:12:11 -05:00
Aiden Cline af0ef00109 Update xAI Grok PDF modalities 2026-05-20 18:06:33 -05:00
neo 9d60164243 feat(stepfun-ai): add international StepFun platform
StepFun runs two separate platforms with distinct accounts/keys:
platform.stepfun.com (China, already covered by providers/stepfun) and
platform.stepfun.ai (international). Keys are not interchangeable
across the two — .ai keys are rejected by api.stepfun.com as
invalid_api_key.

Stepfun's own opencode integration guide instructs users to point at
https://api.stepfun.ai/step_plan/v1. This adds providers/stepfun-ai
for that endpoint, symlinking the shared chat models. Follows the
moonshotai / moonshotai-cn pattern.
2026-05-20 20:43:47 +02:00
arnavchachra 9420048dfe fix crof model limits and reasoning flag to match Crof API 2026-05-18 21:51:19 +05:30
arnavchachra 8db6c27634 add crof provider with 21 models 2026-05-18 21:40:11 +05:30
1679 changed files with 20468 additions and 8199 deletions
+2
View File
@@ -10,6 +10,7 @@ concurrency: ${{ github.workflow }}-${{ github.ref }}
jobs:
deploy:
if: github.repository == 'anomalyco/models.dev'
runs-on: ubuntu-latest
steps:
- name: Checkout code
@@ -35,3 +36,4 @@ jobs:
- run: bun sst deploy --stage=dev
env:
CLOUDFLARE_API_TOKEN: ${{ secrets.CLOUDFLARE_API_TOKEN }}
CLOUDFLARE_DEFAULT_ACCOUNT_ID: ${{ secrets.CLOUDFLARE_DEFAULT_ACCOUNT_ID }}
+38 -15
View File
@@ -2,7 +2,7 @@ name: Sync Model Catalogs
on:
schedule:
- cron: "17 8 * * *"
- cron: "17 * * * *"
workflow_dispatch:
permissions:
@@ -13,20 +13,38 @@ permissions:
concurrency: ${{ github.workflow }}-${{ github.ref }}
jobs:
providers:
if: github.repository == 'anomalyco/models.dev'
runs-on: ubuntu-latest
outputs:
matrix: ${{ steps.providers.outputs.matrix }}
steps:
- name: Checkout code
uses: actions/checkout@34e114876b0b11c390a56381ad16ebd13914f8d5
with:
ref: dev
- name: Setup Bun
uses: oven-sh/setup-bun@f4d14e03ff726c06358e5557344e1da148b56cf7
with:
bun-version: latest
- name: Install dependencies
run: bun install
- name: List sync providers
id: providers
run: |
matrix="$(bun models:sync --list-providers)"
echo "matrix=$matrix" >> "$GITHUB_OUTPUT"
sync:
needs: providers
if: github.repository == 'anomalyco/models.dev'
runs-on: ubuntu-latest
strategy:
fail-fast: false
matrix:
include:
- group: aggregators
title: "chore(sync): update aggregator model catalogs"
branch: automation/sync-models-aggregators
labels: automation,model-sync,sync-group:aggregators,provider:openrouter
- group: cloudflare
title: "chore(sync): update cloudflare model catalogs"
branch: automation/sync-models-cloudflare
labels: automation,model-sync,sync-group:cloudflare,provider:cloudflare-workers-ai
matrix: ${{ fromJSON(needs.providers.outputs.matrix) }}
steps:
- name: Checkout code
@@ -43,9 +61,13 @@ jobs:
run: bun install
- name: Sync model catalogs
run: bun models:sync ${{ matrix.group }}
run: bun models:sync ${{ matrix.provider }}
env:
OPENROUTER_API_KEY: ${{ secrets.OPENROUTER_API_KEY }}
GOOGLE_API_KEY: ${{ secrets.GOOGLE_API_KEY }}
GEMINI_API_KEY: ${{ secrets.GEMINI_API_KEY }}
GOOGLE_GENERATIVE_AI_API_KEY: ${{ secrets.GOOGLE_GENERATIVE_AI_API_KEY }}
XAI_API_KEY: ${{ secrets.XAI_API_KEY }}
CLOUDFLARE_WORKERS_AI_SYNC_ACCOUNT_ID: ${{ secrets.CLOUDFLARE_WORKERS_AI_SYNC_ACCOUNT_ID }}
CLOUDFLARE_WORKERS_AI_SYNC_API_TOKEN: ${{ secrets.CLOUDFLARE_WORKERS_AI_SYNC_API_TOKEN }}
@@ -55,9 +77,9 @@ jobs:
- name: Create pull request
env:
GH_TOKEN: ${{ github.token }}
BRANCH: ${{ matrix.branch }}
LABELS: ${{ matrix.labels }}
TITLE: ${{ matrix.title }}
BRANCH: automation/sync-models-${{ matrix.provider }}
LABELS: automation,model-sync,provider:${{ matrix.provider }}
TITLE: "chore(sync): update ${{ matrix.name }} model catalog"
run: |
if [ -z "$(git status --porcelain -- providers)" ]; then
echo "No model catalog changes found."
@@ -66,6 +88,7 @@ jobs:
git config user.name "github-actions[bot]"
git config user.email "41898282+github-actions[bot]@users.noreply.github.com"
git fetch --no-tags --depth=1 origin "+refs/heads/$BRANCH:refs/remotes/origin/$BRANCH" || true
git checkout -B "$BRANCH"
git add providers
git commit -m "$TITLE"
+380
View File
@@ -0,0 +1,380 @@
{
"name": ".opencode",
"lockfileVersion": 3,
"requires": true,
"packages": {
"": {
"dependencies": {
"@opencode-ai/plugin": "1.15.13"
}
},
"node_modules/@msgpackr-extract/msgpackr-extract-darwin-arm64": {
"version": "3.0.4",
"resolved": "https://registry.npmjs.org/@msgpackr-extract/msgpackr-extract-darwin-arm64/-/msgpackr-extract-darwin-arm64-3.0.4.tgz",
"integrity": "sha512-LCkGo6JDfaBhgST7UpPWgNgLINpcpabaHfyz5OBx75nUYxBsaEPxjnyNjWpeb/xBup/682QnBfRBy2/LvPutZQ==",
"cpu": [
"arm64"
],
"license": "MIT",
"optional": true,
"os": [
"darwin"
]
},
"node_modules/@msgpackr-extract/msgpackr-extract-darwin-x64": {
"version": "3.0.4",
"resolved": "https://registry.npmjs.org/@msgpackr-extract/msgpackr-extract-darwin-x64/-/msgpackr-extract-darwin-x64-3.0.4.tgz",
"integrity": "sha512-zExlW9zUJKZH/tOtVMttwjKa4Xm/3KcNjnE3dPN92uCktwavMxpgCA3MoJK/DOnTWsQgo224OaST27/mPNAf+w==",
"cpu": [
"x64"
],
"license": "MIT",
"optional": true,
"os": [
"darwin"
]
},
"node_modules/@msgpackr-extract/msgpackr-extract-linux-arm": {
"version": "3.0.4",
"resolved": "https://registry.npmjs.org/@msgpackr-extract/msgpackr-extract-linux-arm/-/msgpackr-extract-linux-arm-3.0.4.tgz",
"integrity": "sha512-Tg3yX65f5GbtXLkrYEHE5oibZG9epyYWas7FogTTEJeDEF9JlXJzKgXaNhT3UXlTOeA+AfZpYZYZ0uPj7Cfquw==",
"cpu": [
"arm"
],
"license": "MIT",
"optional": true,
"os": [
"linux"
]
},
"node_modules/@msgpackr-extract/msgpackr-extract-linux-arm64": {
"version": "3.0.4",
"resolved": "https://registry.npmjs.org/@msgpackr-extract/msgpackr-extract-linux-arm64/-/msgpackr-extract-linux-arm64-3.0.4.tgz",
"integrity": "sha512-dgX0P/9wGPJeHFBG+ZmhgE6bmtMt7NP5CRBGyyktpopdk/mW4POnrpQsSLtKI1dwpc+pPLuXHDh6vvskyQE/sw==",
"cpu": [
"arm64"
],
"license": "MIT",
"optional": true,
"os": [
"linux"
]
},
"node_modules/@msgpackr-extract/msgpackr-extract-linux-x64": {
"version": "3.0.4",
"resolved": "https://registry.npmjs.org/@msgpackr-extract/msgpackr-extract-linux-x64/-/msgpackr-extract-linux-x64-3.0.4.tgz",
"integrity": "sha512-8TNXMEjJc3QEy7R/x1INhgiU+XakDAFUzBhaz7+Rbrs8NH5UQeHQxxmzsSBJGyV6I1jW79undiQm8tOI+D+8FQ==",
"cpu": [
"x64"
],
"license": "MIT",
"optional": true,
"os": [
"linux"
]
},
"node_modules/@msgpackr-extract/msgpackr-extract-win32-x64": {
"version": "3.0.4",
"resolved": "https://registry.npmjs.org/@msgpackr-extract/msgpackr-extract-win32-x64/-/msgpackr-extract-win32-x64-3.0.4.tgz",
"integrity": "sha512-CmCXPQrkbwExx3j946/PtHWHbYJiCRBRDl4BlkRQcJB/YOwQxJRTpoo7aTsortjgoJ1x7opzTSxn7C+ASSLVjQ==",
"cpu": [
"x64"
],
"license": "MIT",
"optional": true,
"os": [
"win32"
]
},
"node_modules/@opencode-ai/plugin": {
"version": "1.15.13",
"resolved": "https://registry.npmjs.org/@opencode-ai/plugin/-/plugin-1.15.13.tgz",
"integrity": "sha512-NFwZGhmxIPijtfz9swPJXDmhOpq4UWP8WjEE7GEMr7FwtJrK/hv6v36nFimed5+OKk+pQCrTJn/vhRW7Io72IA==",
"license": "MIT",
"dependencies": {
"@opencode-ai/sdk": "1.15.13",
"effect": "4.0.0-beta.66",
"zod": "4.1.8"
},
"peerDependencies": {
"@opentui/core": ">=0.2.16",
"@opentui/keymap": ">=0.2.16",
"@opentui/solid": ">=0.2.16"
},
"peerDependenciesMeta": {
"@opentui/core": {
"optional": true
},
"@opentui/keymap": {
"optional": true
},
"@opentui/solid": {
"optional": true
}
}
},
"node_modules/@opencode-ai/sdk": {
"version": "1.15.13",
"resolved": "https://registry.npmjs.org/@opencode-ai/sdk/-/sdk-1.15.13.tgz",
"integrity": "sha512-4TwojIoQ8EG6/mVBuUVYZXiFcwNmiiytEnjnvyuvSJjGwFIlw2YIBFxtSVC3FbwwbwHT63teh1RHiQUUC4U5xw==",
"license": "MIT",
"dependencies": {
"cross-spawn": "7.0.6"
}
},
"node_modules/@standard-schema/spec": {
"version": "1.1.0",
"resolved": "https://registry.npmjs.org/@standard-schema/spec/-/spec-1.1.0.tgz",
"integrity": "sha512-l2aFy5jALhniG5HgqrD6jXLi/rUWrKvqN/qJx6yoJsgKhblVd+iqqU4RCXavm/jPityDo5TCvKMnpjKnOriy0w==",
"license": "MIT"
},
"node_modules/cross-spawn": {
"version": "7.0.6",
"resolved": "https://registry.npmjs.org/cross-spawn/-/cross-spawn-7.0.6.tgz",
"integrity": "sha512-uV2QOWP2nWzsy2aMp8aRibhi9dlzF5Hgh5SHaB9OiTGEyDTiJJyx0uy51QXdyWbtAHNua4XJzUKca3OzKUd3vA==",
"license": "MIT",
"dependencies": {
"path-key": "^3.1.0",
"shebang-command": "^2.0.0",
"which": "^2.0.1"
},
"engines": {
"node": ">= 8"
}
},
"node_modules/detect-libc": {
"version": "2.1.2",
"resolved": "https://registry.npmjs.org/detect-libc/-/detect-libc-2.1.2.tgz",
"integrity": "sha512-Btj2BOOO83o3WyH59e8MgXsxEQVcarkUOpEYrubB0urwnN10yQ364rsiByU11nZlqWYZm05i/of7io4mzihBtQ==",
"license": "Apache-2.0",
"optional": true,
"engines": {
"node": ">=8"
}
},
"node_modules/effect": {
"version": "4.0.0-beta.66",
"resolved": "https://registry.npmjs.org/effect/-/effect-4.0.0-beta.66.tgz",
"integrity": "sha512-4arEr62cziFa8BBVDUwJCJJmaVepXf/kRg7KtC0h8+bufngscrHbwWFhr9c+HonwOF+31U3iD3xUJmw9KzX7Dw==",
"license": "MIT",
"dependencies": {
"@standard-schema/spec": "^1.1.0",
"fast-check": "^4.6.0",
"find-my-way-ts": "^0.1.6",
"ini": "^6.0.0",
"kubernetes-types": "^1.30.0",
"msgpackr": "^1.11.9",
"multipasta": "^0.2.7",
"toml": "^4.1.1",
"uuid": "^13.0.0",
"yaml": "^2.8.3"
}
},
"node_modules/fast-check": {
"version": "4.8.0",
"resolved": "https://registry.npmjs.org/fast-check/-/fast-check-4.8.0.tgz",
"integrity": "sha512-GOJ158CUMnN6cSahsv4+ExARvIDuzzinFjkp0E9WtiBa5zcVeLozVkWaE4IzFcc+Y48Wp1EDlUZsXRyAztQcSg==",
"funding": [
{
"type": "individual",
"url": "https://github.com/sponsors/dubzzz"
},
{
"type": "opencollective",
"url": "https://opencollective.com/fast-check"
}
],
"license": "MIT",
"dependencies": {
"pure-rand": "^8.0.0"
},
"engines": {
"node": ">=12.17.0"
}
},
"node_modules/find-my-way-ts": {
"version": "0.1.6",
"resolved": "https://registry.npmjs.org/find-my-way-ts/-/find-my-way-ts-0.1.6.tgz",
"integrity": "sha512-a85L9ZoXtNAey3Y6Z+eBWW658kO/MwR7zIafkIUPUMf3isZG0NCs2pjW2wtjxAKuJPxMAsHUIP4ZPGv0o5gyTA==",
"license": "MIT"
},
"node_modules/ini": {
"version": "6.0.0",
"resolved": "https://registry.npmjs.org/ini/-/ini-6.0.0.tgz",
"integrity": "sha512-IBTdIkzZNOpqm7q3dRqJvMaldXjDHWkEDfrwGEQTs5eaQMWV+djAhR+wahyNNMAa+qpbDUhBMVt4ZKNwpPm7xQ==",
"license": "ISC",
"engines": {
"node": "^20.17.0 || >=22.9.0"
}
},
"node_modules/isexe": {
"version": "2.0.0",
"resolved": "https://registry.npmjs.org/isexe/-/isexe-2.0.0.tgz",
"integrity": "sha512-RHxMLp9lnKHGHRng9QFhRCMbYAcVpn69smSGcq3f36xjgVVWThj4qqLbTLlq7Ssj8B+fIQ1EuCEGI2lKsyQeIw==",
"license": "ISC"
},
"node_modules/kubernetes-types": {
"version": "1.30.0",
"resolved": "https://registry.npmjs.org/kubernetes-types/-/kubernetes-types-1.30.0.tgz",
"integrity": "sha512-Dew1okvhM/SQcIa2rcgujNndZwU8VnSapDgdxlYoB84ZlpAD43U6KLAFqYo17ykSFGHNPrg0qry0bP+GJd9v7Q==",
"license": "Apache-2.0"
},
"node_modules/msgpackr": {
"version": "1.11.12",
"resolved": "https://registry.npmjs.org/msgpackr/-/msgpackr-1.11.12.tgz",
"integrity": "sha512-RBdJ1Un7yGlXWajrkxcSa93nvQ0w4zBf60c0yYv7YtBelP8H2FA7XsfBbMHtXKXUMUxH7zV3Zuozh+kUQWhHvg==",
"license": "MIT",
"optionalDependencies": {
"msgpackr-extract": "^3.0.2"
}
},
"node_modules/msgpackr-extract": {
"version": "3.0.4",
"resolved": "https://registry.npmjs.org/msgpackr-extract/-/msgpackr-extract-3.0.4.tgz",
"integrity": "sha512-4kmO/MdyUIkLIvTPr8VHLil4AtoKIoniWPIEk5+CDy0xnWC84azhSFmuJ7PxZdsYtiP5kEeQsORAVIeMgxT+Hw==",
"hasInstallScript": true,
"license": "MIT",
"optional": true,
"dependencies": {
"node-gyp-build-optional-packages": "5.2.2"
},
"bin": {
"download-msgpackr-prebuilds": "bin/download-prebuilds.js"
},
"optionalDependencies": {
"@msgpackr-extract/msgpackr-extract-darwin-arm64": "3.0.4",
"@msgpackr-extract/msgpackr-extract-darwin-x64": "3.0.4",
"@msgpackr-extract/msgpackr-extract-linux-arm": "3.0.4",
"@msgpackr-extract/msgpackr-extract-linux-arm64": "3.0.4",
"@msgpackr-extract/msgpackr-extract-linux-x64": "3.0.4",
"@msgpackr-extract/msgpackr-extract-win32-x64": "3.0.4"
}
},
"node_modules/multipasta": {
"version": "0.2.7",
"resolved": "https://registry.npmjs.org/multipasta/-/multipasta-0.2.7.tgz",
"integrity": "sha512-KPA58d68KgGil15oDqXjkUBEBYc00XvbPj5/X+dyzeo/lWm9Nc25pQRlf1D+gv4OpK7NM0J1odrbu9JNNGvynA==",
"license": "MIT"
},
"node_modules/node-gyp-build-optional-packages": {
"version": "5.2.2",
"resolved": "https://registry.npmjs.org/node-gyp-build-optional-packages/-/node-gyp-build-optional-packages-5.2.2.tgz",
"integrity": "sha512-s+w+rBWnpTMwSFbaE0UXsRlg7hU4FjekKU4eyAih5T8nJuNZT1nNsskXpxmeqSK9UzkBl6UgRlnKc8hz8IEqOw==",
"license": "MIT",
"optional": true,
"dependencies": {
"detect-libc": "^2.0.1"
},
"bin": {
"node-gyp-build-optional-packages": "bin.js",
"node-gyp-build-optional-packages-optional": "optional.js",
"node-gyp-build-optional-packages-test": "build-test.js"
}
},
"node_modules/path-key": {
"version": "3.1.1",
"resolved": "https://registry.npmjs.org/path-key/-/path-key-3.1.1.tgz",
"integrity": "sha512-ojmeN0qd+y0jszEtoY48r0Peq5dwMEkIlCOu6Q5f41lfkswXuKtYrhgoTpLnyIcHm24Uhqx+5Tqm2InSwLhE6Q==",
"license": "MIT",
"engines": {
"node": ">=8"
}
},
"node_modules/pure-rand": {
"version": "8.4.0",
"resolved": "https://registry.npmjs.org/pure-rand/-/pure-rand-8.4.0.tgz",
"integrity": "sha512-IoM8YF/jY0hiugFo/wOWqfmarlE6J0wc6fDK1PhftMk7MGhVZl88sZimmqBBFomLOCSmcCCpsfj7wXASCpvK9A==",
"funding": [
{
"type": "individual",
"url": "https://github.com/sponsors/dubzzz"
},
{
"type": "opencollective",
"url": "https://opencollective.com/fast-check"
}
],
"license": "MIT"
},
"node_modules/shebang-command": {
"version": "2.0.0",
"resolved": "https://registry.npmjs.org/shebang-command/-/shebang-command-2.0.0.tgz",
"integrity": "sha512-kHxr2zZpYtdmrN1qDjrrX/Z1rR1kG8Dx+gkpK1G4eXmvXswmcE1hTWBWYUzlraYw1/yZp6YuDY77YtvbN0dmDA==",
"license": "MIT",
"dependencies": {
"shebang-regex": "^3.0.0"
},
"engines": {
"node": ">=8"
}
},
"node_modules/shebang-regex": {
"version": "3.0.0",
"resolved": "https://registry.npmjs.org/shebang-regex/-/shebang-regex-3.0.0.tgz",
"integrity": "sha512-7++dFhtcx3353uBaq8DDR4NuxBetBzC7ZQOhmTQInHEd6bSrXdiEyzCvG07Z44UYdLShWUyXt5M/yhz8ekcb1A==",
"license": "MIT",
"engines": {
"node": ">=8"
}
},
"node_modules/toml": {
"version": "4.1.1",
"resolved": "https://registry.npmjs.org/toml/-/toml-4.1.1.tgz",
"integrity": "sha512-EBJnVBr3dTXdA89WVFoAIPUqkBjxPMwRqsfuo1r240tKFHXv3zgca4+NJib/h6TyvGF7vOawz0jGuryJCdNHrw==",
"license": "MIT",
"engines": {
"node": ">=20"
}
},
"node_modules/uuid": {
"version": "13.0.2",
"resolved": "https://registry.npmjs.org/uuid/-/uuid-13.0.2.tgz",
"integrity": "sha512-vzi9uRZ926x4XV73S/4qQaTwPXM2JBj6/6lI/byHH1jOpCzb0zDbfytgA9LcN/hzb2l7WQSQnxITOVx5un/wGw==",
"funding": [
"https://github.com/sponsors/broofa",
"https://github.com/sponsors/ctavan"
],
"license": "MIT",
"bin": {
"uuid": "dist-node/bin/uuid"
}
},
"node_modules/which": {
"version": "2.0.2",
"resolved": "https://registry.npmjs.org/which/-/which-2.0.2.tgz",
"integrity": "sha512-BLI3Tl1TW3Pvl70l3yq3Y64i+awpwXqsGBYWkkqMtnbXgrMD+yj7rhW0kuEDxzJaYXGjEW5ogapKNMEKNMjibA==",
"license": "ISC",
"dependencies": {
"isexe": "^2.0.0"
},
"bin": {
"node-which": "bin/node-which"
},
"engines": {
"node": ">= 8"
}
},
"node_modules/yaml": {
"version": "2.9.0",
"resolved": "https://registry.npmjs.org/yaml/-/yaml-2.9.0.tgz",
"integrity": "sha512-2AvhNX3mb8zd6Zy7INTtSpl1F15HW6Wnqj0srWlkKLcpYl/gMIMJiyuGq2KeI2YFxUPjdlB+3Lc10seMLtL4cA==",
"license": "ISC",
"bin": {
"yaml": "bin.mjs"
},
"engines": {
"node": ">= 14.6"
},
"funding": {
"url": "https://github.com/sponsors/eemeli"
}
},
"node_modules/zod": {
"version": "4.1.8",
"license": "MIT",
"funding": {
"url": "https://github.com/sponsors/colinhacks"
}
}
}
}
+14 -15
View File
@@ -31,26 +31,25 @@
## Model Configuration
- Model `id` is **auto-injected** from filename (minus `.toml`) — never put `id` in TOML files
- Models may reuse another model's definition via `extends` (see below); otherwise the full definition must be present in the file
- Provider models may reuse provider-agnostic facts from `models/` via `base_model`; otherwise the full provider model definition must be present in the file
- Schema uses `.strict()` — extra fields cause validation errors
### `[extends]` (inheritance between models)
- Syntax — a table at the top of the TOML:
### Model metadata and `base_model`
- Provider-agnostic model facts live under `models/<provider>/<model>.toml`
- Provider TOMLs can inherit those facts with:
```toml
[extends]
from = "<provider-id>/<model-id>" # required
omit = ["experimental.modes.fast"] # optional, dot-path strings
base_model = "<provider-id>/<model-id>"
base_model_omit = ["limit.input"] # optional, dot-path strings
```
Example: `from = "anthropic/claude-opus-4-6"`
- Resolved at parse time in `generate()`; the final JSON output contains **no** `extends` field — it exists only to cut duplication in the TOMLs
Example: `base_model = "anthropic/claude-opus-4-6"`
- Resolved at parse time in `generate()`; the final provider JSON output contains **no** `base_model` or `base_model_omit` fields
- Merge semantics:
- Plain objects (`[cost]`, `[limit]`, `[modalities]`, `[provider]`, `[experimental]`, …) are **deep-merged**
- Plain objects from metadata and provider TOML (`[limit]`, `[modalities]`, …) are **deep-merged**
- Arrays (e.g. `modalities.input`) and primitives are **replaced** wholesale by the child
- Any field the child omits is inherited verbatim from the base
- `omit` runs **after** the merge and deletes each dot-path from the result (used when the child needs to *remove* something the base defines, e.g. a provider-specific experimental mode). Every listed path must exist in the merged model, else an error is thrown. Ancestor tables that become empty as a result are also pruned, so `omit = ["experimental.modes.fast"]` yields no `experimental` key in the final JSON when `fast` was the only mode.
- Chains are allowed (A extends B extends C); cycles throw
- The base model must exist; `[extends.from]` pointing at a missing provider/model is an error
- The `extends` table is stripped before schema validation, so the merged result must still satisfy the strict `Model` schema
- Any provider field omitted is inherited verbatim from model metadata
- `cost`, `provider`, `experimental`, `reasoning_options`, `interleaved`, and `status` are provider-specific and must be declared in provider TOMLs when needed
- `base_model_omit` runs **after** the merge and deletes each dot-path from the result. Missing paths are ignored. Ancestor tables that become empty as a result are also pruned.
- The base model metadata file must exist; `base_model` pointing at a missing `models/` entry is an error
### Bedrock Naming Patterns
- Dated models: `-v1:0` suffix (`anthropic.claude-3-5-sonnet-20241022-v1:0.toml`)
@@ -72,4 +71,4 @@
| `attachment`, `reasoning`, `tool_call`, `open_weights` | Yes | Boolean capabilities |
| `cost`, `limit`, `modalities` | Yes | Objects with their own required fields |
| `family`, `knowledge`, `temperature`, `structured_output` | No | Optional metadata |
| `status` | No | Use for `"alpha"`, `"beta"`, `"deprecated"` lifecycle |
| `status` | No | Use for `"alpha"`, `"beta"`, `"deprecated"` lifecycle |
+90 -13
View File
@@ -24,6 +24,18 @@ curl https://models.dev/api.json
Use the **Model ID** field to do a lookup on any model; it's the identifier used by [AI SDK](https://ai-sdk.dev/).
Provider-agnostic model metadata is available separately:
```bash
curl https://models.dev/models.json
```
Use this for facts about the model itself, independent of where it is served. If you need both provider endpoints and model-only metadata in one response:
```bash
curl https://models.dev/catalog.json
```
### Logos
Provider logos are available as SVG files:
@@ -40,7 +52,71 @@ The data is stored in the repo as TOML files; organized by provider and model. T
We need your help keeping the data up to date.
### Adding a New Model
### Adding Model Metadata
Model-only facts live in `models/`, using the same path-style IDs as provider models. For example, `models/openai/gpt-5.toml` defines metadata for the underlying GPT-5 model, while `providers/openai/models/gpt-5.toml` defines OpenAI-specific serving details such as pricing.
Use model metadata for provider-agnostic facts:
- `name`, `family`, `release_date`, `last_updated`, `knowledge`
- `attachment`, `reasoning`, `tool_call`, `structured_output`, `temperature`
- `[limit]` defaults like context, input, and output token limits
- `[modalities]` defaults
- `open_weights`, `license`, `links`, `weights`, and `benchmarks`
Example:
```toml
name = "GPT-5"
family = "gpt"
release_date = "2025-08-07"
last_updated = "2025-08-07"
attachment = true
reasoning = true
temperature = false
tool_call = true
structured_output = true
open_weights = false
[limit]
context = 400_000
input = 272_000
output = 128_000
[modalities]
input = ["text", "image"]
output = ["text"]
[[benchmarks]]
name = "Benchmark Name"
score = 72.5
metric = "accuracy"
source = "https://example.com/results"
[[weights]]
label = "Model weights"
url = "https://huggingface.co/example/model"
format = "safetensors"
```
Provider TOMLs can inherit these facts with `base_model` and then keep only provider-specific fields or overrides:
```toml
base_model = "openai/gpt-5"
[cost]
input = 1.25
output = 10.00
cache_read = 0.125
[limit]
context = 200_000 # optional provider override
output = 32_000
```
Provider fields win over model metadata during generation. Use this when the underlying model is the same but a provider serves it with different context limits, modalities, features, or pricing.
### Adding a New Provider Model
To add a new model, start by checking if the provider already exists in the `providers/` directory. If not, then:
@@ -120,30 +196,31 @@ output = ["text"] # Supported output modalities
field = "reasoning_content" # Name of the interleaved field "reasoning_content" or "reasoning_details"
```
#### 3a. Reuse an Existing Model with `extends`
#### 3a. Reuse Model Metadata with `base_model`
For wrapper providers that mirror a model from another provider, prefer reusing the canonical model definition instead of duplicating the whole file.
For wrapper providers that mirror an existing model, prefer referencing the model-only metadata instead of duplicating provider-agnostic fields.
Use `extends` only for non-first-party wrappers and mirrors. Do not use it inside the actual lab provider directories that act as the canonical source for a model family, for example `providers/anthropic/`, `providers/openai/`, `providers/google/`, `providers/xai/`, `providers/minimax/`, or `providers/moonshot/`.
Use `base_model` when the provider serves the same underlying model and only provider-specific fields differ.
```toml
[extends]
from = "anthropic/claude-opus-4-6"
omit = ["experimental.modes.fast"]
base_model = "anthropic/claude-opus-4-6"
[provider]
npm = "@ai-sdk/anthropic"
[cost]
input = 5.00
output = 25.00
```
Rules:
- `from` must point to another model using `<provider>/<model-id>`.
- `omit` is optional and removes fields after the inherited model and local overrides are merged.
- `base_model` must point to a TOML file in `models/` using `<provider>/<model-id>`.
- You can override any top-level model field locally.
- If you override a nested table like `[cost]`, `[limit]`, or `[modalities]`, include the full values needed for that table.
- `base_model_omit` is optional and removes inherited model metadata fields after local overrides are merged. Use dot-path strings, for example `base_model_omit = ["limit.input"]`.
- `id` still comes from the filename; do not add it to the TOML.
Use `extends` when the wrapper model is materially the same as the source model and only differs by a small set of overrides or omitted fields.
Use `base_model` when the wrapper model is materially the same as the source model and only differs by provider-specific pricing, limits, modalities, provider request shape, or lifecycle flags.
Sync and generator scripts should preserve existing `base_model` / `base_model_omit` fields when updating provider TOMLs. Do not use legacy `[extends]` tables.
#### 4. Submit a Pull Request
@@ -161,7 +238,7 @@ There's a GitHub Action that will automatically validate your submission against
- Values are within acceptable ranges
- TOML syntax is valid
When converting existing wrapper models to `extends`, compare generated output before and after the change:
When moving existing provider fields into model metadata, compare generated output before and after the change:
```bash
bun run compare:migrations
+18
View File
@@ -0,0 +1,18 @@
name = "Qwen Flash"
family = "qwen"
release_date = "2025-07-28"
last_updated = "2025-07-28"
attachment = false
reasoning = true
temperature = true
tool_call = true
knowledge = "2024-04"
open_weights = false
[limit]
context = 1_000_000
output = 32_768
[modalities]
input = ["text"]
output = ["text"]
+25
View File
@@ -0,0 +1,25 @@
name = "Qwen Max"
family = "qwen"
release_date = "2024-04-03"
last_updated = "2025-01-25"
attachment = false
reasoning = false
temperature = true
tool_call = true
knowledge = "2024-04"
open_weights = false
[limit]
context = 32_768
output = 8_192
[modalities]
input = ["text"]
output = ["text"]
[[benchmarks]]
name = "Aider Polyglot"
score = 21.8
metric = "percent correct"
source = "https://aider.chat/docs/leaderboards/"
date = "2025-01-28"
+18
View File
@@ -0,0 +1,18 @@
name = "Qwen-Omni Turbo"
family = "qwen"
release_date = "2025-01-19"
last_updated = "2025-03-26"
attachment = false
reasoning = false
temperature = true
tool_call = true
knowledge = "2024-04"
open_weights = false
[limit]
context = 32_768
output = 2_048
[modalities]
input = ["text", "image", "audio", "video"]
output = ["text", "audio"]
+18
View File
@@ -0,0 +1,18 @@
name = "Qwen Plus"
family = "qwen"
release_date = "2024-01-25"
last_updated = "2025-09-11"
attachment = false
reasoning = true
temperature = true
tool_call = true
knowledge = "2024-04"
open_weights = false
[limit]
context = 1_000_000
output = 32_768
[modalities]
input = ["text"]
output = ["text"]
+18
View File
@@ -0,0 +1,18 @@
name = "Qwen Turbo"
family = "qwen"
release_date = "2024-11-01"
last_updated = "2025-04-28"
attachment = false
reasoning = true
temperature = true
tool_call = true
knowledge = "2024-04"
open_weights = false
[limit]
context = 1_000_000
output = 16_384
[modalities]
input = ["text"]
output = ["text"]
+18
View File
@@ -0,0 +1,18 @@
name = "Qwen-VL Max"
family = "qwen"
release_date = "2024-04-08"
last_updated = "2025-08-13"
attachment = false
reasoning = false
temperature = true
tool_call = true
knowledge = "2024-04"
open_weights = false
[limit]
context = 131_072
output = 8_192
[modalities]
input = ["text", "image"]
output = ["text"]
+18
View File
@@ -0,0 +1,18 @@
name = "Qwen-VL Plus"
family = "qwen"
release_date = "2024-01-25"
last_updated = "2025-08-15"
attachment = false
reasoning = false
temperature = true
tool_call = true
knowledge = "2024-04"
open_weights = false
[limit]
context = 131_072
output = 8_192
[modalities]
input = ["text", "image"]
output = ["text"]
@@ -0,0 +1,22 @@
name = "Qwen2.5-VL 72B Instruct"
family = "qwen"
release_date = "2024-09"
last_updated = "2024-09"
attachment = false
reasoning = false
temperature = true
tool_call = true
knowledge = "2024-04"
open_weights = true
[limit]
context = 131_072
output = 8_192
[modalities]
input = ["text", "image"]
output = ["text"]
[[weights]]
label = "Hugging Face"
url = "https://huggingface.co/Qwen/Qwen2.5-VL-72B-Instruct"
+36
View File
@@ -0,0 +1,36 @@
name = "Qwen3 235B-A22B"
family = "qwen"
release_date = "2025-04"
last_updated = "2025-04"
attachment = false
reasoning = true
temperature = true
tool_call = true
knowledge = "2025-04"
open_weights = true
[limit]
context = 131_072
output = 16_384
[modalities]
input = ["text"]
output = ["text"]
[[weights]]
label = "Hugging Face"
url = "https://huggingface.co/Qwen/Qwen3-235B-A22B"
[[benchmarks]]
name = "Aider Polyglot"
score = 59.6
metric = "percent correct"
source = "https://aider.chat/docs/leaderboards/"
date = "2025-05-09"
[[benchmarks]]
name = "SWE-Bench Pro"
score = 21.41
metric = "resolve rate"
dataset = "public"
source = "https://labs.scale.com/leaderboard/swe_bench_pro_public"
+29
View File
@@ -0,0 +1,29 @@
name = "Qwen3 32B"
family = "qwen"
release_date = "2025-04"
last_updated = "2025-04"
attachment = false
reasoning = true
temperature = true
tool_call = true
knowledge = "2025-04"
open_weights = true
[limit]
context = 131_072
output = 16_384
[modalities]
input = ["text"]
output = ["text"]
[[weights]]
label = "Hugging Face"
url = "https://huggingface.co/Qwen/Qwen3-32B"
[[benchmarks]]
name = "Aider Polyglot"
score = 40.0
metric = "percent correct"
source = "https://aider.chat/docs/leaderboards/"
date = "2025-05-08"
@@ -0,0 +1,43 @@
name = "Qwen3-Coder 30B-A3B Instruct"
family = "qwen"
release_date = "2025-04"
last_updated = "2025-04"
attachment = false
reasoning = false
temperature = true
tool_call = true
knowledge = "2025-04"
open_weights = true
[limit]
context = 262_144
output = 65_536
[modalities]
input = ["text"]
output = ["text"]
[[weights]]
label = "Hugging Face"
url = "https://huggingface.co/Qwen/Qwen3-Coder-30B-A3B-Instruct"
[[benchmarks]]
name = "Artificial Analysis Coding Index"
score = 19.4
metric = "index"
source = "https://openrouter.ai/qwen/qwen3-coder-30b-a3b-instruct/benchmarks"
date = "2026-06-02"
[[benchmarks]]
name = "SciCode"
score = 27.8
metric = "percent correct"
source = "https://openrouter.ai/qwen/qwen3-coder-30b-a3b-instruct/benchmarks"
date = "2026-06-02"
[[benchmarks]]
name = "Terminal-Bench Hard"
score = 15.2
metric = "success rate"
source = "https://openrouter.ai/qwen/qwen3-coder-30b-a3b-instruct/benchmarks"
date = "2026-06-02"
@@ -0,0 +1,29 @@
name = "Qwen3-Coder 480B-A35B Instruct"
family = "qwen"
release_date = "2025-04"
last_updated = "2025-04"
attachment = false
reasoning = false
temperature = true
tool_call = true
knowledge = "2025-04"
open_weights = true
[limit]
context = 262_144
output = 65_536
[modalities]
input = ["text"]
output = ["text"]
[[weights]]
label = "Hugging Face"
url = "https://huggingface.co/Qwen/Qwen3-Coder-480B-A35B-Instruct"
[[benchmarks]]
name = "SWE-Bench Pro"
score = 38.7
metric = "resolve rate"
dataset = "public"
source = "https://labs.scale.com/leaderboard/swe_bench_pro_public"
@@ -1,19 +1,16 @@
name = "GLM 4.5 Air Derestricted Steam"
name = "Qwen3 Coder Flash"
family = "qwen"
release_date = "2025-07-28"
last_updated = "2025-07-28"
attachment = false
reasoning = false
tool_call = false
structured_output = false
temperature = true
tool_call = true
knowledge = "2025-04"
open_weights = false
[cost]
input = 0.306
output = 0.306
[limit]
context = 220_600
input = 220_600
context = 1_000_000
output = 65_536
[modalities]
@@ -1,21 +1,17 @@
name = "Qwen3 Coder 480B TEE"
name = "Qwen3 Coder Plus"
family = "qwen"
release_date = "2025-07-23"
last_updated = "2025-07-23"
attachment = false
reasoning = false
tool_call = false
structured_output = false
temperature = true
tool_call = true
knowledge = "2025-04"
open_weights = false
[cost]
input = 1.5
output = 2.00
[limit]
context = 128_000
input = 128_000
output = 32_768
context = 1_048_576
output = 65_536
[modalities]
input = ["text"]
+39
View File
@@ -0,0 +1,39 @@
name = "Qwen3 Max"
family = "qwen"
release_date = "2025-09-23"
last_updated = "2025-09-23"
attachment = false
reasoning = false
temperature = true
tool_call = true
knowledge = "2025-04"
open_weights = false
[limit]
context = 262_144
output = 65_536
[modalities]
input = ["text"]
output = ["text"]
[[benchmarks]]
name = "Artificial Analysis Coding Index"
score = 26.4
metric = "index"
source = "https://openrouter.ai/qwen/qwen3-max/benchmarks"
date = "2026-05-30"
[[benchmarks]]
name = "SciCode"
score = 38.3
metric = "percent correct"
source = "https://openrouter.ai/qwen/qwen3-max/benchmarks"
date = "2026-05-30"
[[benchmarks]]
name = "Terminal-Bench Hard"
score = 20.5
metric = "success rate"
source = "https://openrouter.ai/qwen/qwen3-max/benchmarks"
date = "2026-05-30"
@@ -0,0 +1,22 @@
name = "Qwen3-Next 80B-A3B Instruct"
family = "qwen"
release_date = "2025-09"
last_updated = "2025-09"
attachment = false
reasoning = false
temperature = true
tool_call = true
knowledge = "2025-04"
open_weights = true
[limit]
context = 131_072
output = 32_768
[modalities]
input = ["text"]
output = ["text"]
[[weights]]
label = "Hugging Face"
url = "https://huggingface.co/Qwen/Qwen3-Next-80B-A3B-Instruct"
@@ -0,0 +1,22 @@
name = "Qwen3-Next 80B-A3B (Thinking)"
family = "qwen"
release_date = "2025-09"
last_updated = "2025-09"
attachment = false
reasoning = true
temperature = true
tool_call = true
knowledge = "2025-04"
open_weights = true
[limit]
context = 131_072
output = 32_768
[modalities]
input = ["text"]
output = ["text"]
[[weights]]
label = "Hugging Face"
url = "https://huggingface.co/Qwen/Qwen3-Next-80B-A3B-Thinking"
+18
View File
@@ -0,0 +1,18 @@
name = "Qwen3-VL Plus"
family = "qwen"
release_date = "2025-09-23"
last_updated = "2025-09-23"
attachment = false
reasoning = true
temperature = true
tool_call = true
knowledge = "2025-04"
open_weights = false
[limit]
context = 262_144
output = 32_768
[modalities]
input = ["text", "image"]
output = ["text"]
+28
View File
@@ -0,0 +1,28 @@
name = "Qwen3.5 122B-A10B"
family = "qwen"
release_date = "2026-02-23"
last_updated = "2026-02-23"
attachment = true
reasoning = true
temperature = true
tool_call = true
structured_output = true
open_weights = true
[limit]
context = 262_144
output = 65_536
[modalities]
input = ["text", "image", "video", "audio"]
output = ["text"]
[[weights]]
label = "Hugging Face"
url = "https://huggingface.co/Qwen/Qwen3.5-122B-A10B"
[[benchmarks]]
name = "SWE-Bench Verified"
score = 72
metric = "resolved"
source = "https://huggingface.co/Qwen/Qwen3.5-122B-A10B"
+28
View File
@@ -0,0 +1,28 @@
name = "Qwen3.5 27B"
family = "qwen"
release_date = "2026-02-23"
last_updated = "2026-02-23"
attachment = true
reasoning = true
temperature = true
tool_call = true
structured_output = true
open_weights = true
[limit]
context = 262_144
output = 65_536
[modalities]
input = ["text", "image", "video", "audio"]
output = ["text"]
[[weights]]
label = "Hugging Face"
url = "https://huggingface.co/Qwen/Qwen3.5-27B"
[[benchmarks]]
name = "SWE-Bench Verified"
score = 72.4
metric = "resolved"
source = "https://huggingface.co/Qwen/Qwen3.5-27B"
+22
View File
@@ -0,0 +1,22 @@
name = "Qwen3.5 35B-A3B"
family = "qwen"
release_date = "2026-02-23"
last_updated = "2026-02-23"
attachment = true
reasoning = true
temperature = true
tool_call = true
structured_output = true
open_weights = true
[limit]
context = 262_144
output = 65_536
[modalities]
input = ["text", "image", "video", "audio"]
output = ["text"]
[[weights]]
label = "Hugging Face"
url = "https://huggingface.co/Qwen/Qwen3.5-35B-A3B"
+28
View File
@@ -0,0 +1,28 @@
name = "Qwen3.5 397B-A17B"
family = "qwen"
release_date = "2026-02-15"
last_updated = "2026-02-15"
attachment = true
reasoning = true
temperature = true
tool_call = true
structured_output = true
open_weights = true
[limit]
context = 262_144
output = 65_536
[modalities]
input = ["text", "image", "video", "audio"]
output = ["text"]
[[weights]]
label = "Hugging Face"
url = "https://huggingface.co/Qwen/Qwen3.5-397B-A17B"
[[benchmarks]]
name = "SWE-Bench Verified"
score = 76.4
metric = "resolved"
source = "https://huggingface.co/Qwen/Qwen3.5-397B-A17B"
+18
View File
@@ -0,0 +1,18 @@
name = "Qwen3.5 Plus"
family = "qwen"
release_date = "2026-02-16"
last_updated = "2026-02-16"
attachment = false
reasoning = true
temperature = true
tool_call = true
knowledge = "2025-04"
open_weights = false
[limit]
context = 1_000_000
output = 65_536
[modalities]
input = ["text", "image", "video"]
output = ["text"]
+28
View File
@@ -0,0 +1,28 @@
name = "Qwen3.6 27B"
family = "qwen"
release_date = "2026-04-22"
last_updated = "2026-04-22"
attachment = true
reasoning = true
temperature = true
tool_call = true
structured_output = true
open_weights = true
[limit]
context = 262_144
output = 65_536
[modalities]
input = ["text", "image", "video", "audio"]
output = ["text"]
[[weights]]
label = "Hugging Face"
url = "https://huggingface.co/Qwen/Qwen3.6-27B"
[[benchmarks]]
name = "SWE-Bench Verified"
score = 77.2
metric = "resolved"
source = "https://huggingface.co/Qwen/Qwen3.6-27B"
+28
View File
@@ -0,0 +1,28 @@
name = "Qwen3.6 35B-A3B"
family = "qwen"
release_date = "2026-04-17"
last_updated = "2026-04-17"
attachment = true
reasoning = true
temperature = true
tool_call = true
structured_output = true
open_weights = true
[limit]
context = 262_144
output = 65_536
[modalities]
input = ["text", "image", "video", "audio"]
output = ["text"]
[[weights]]
label = "Hugging Face"
url = "https://huggingface.co/Qwen/Qwen3.6-35B-A3B"
[[benchmarks]]
name = "SWE-Bench Verified"
score = 73.4
metric = "resolved"
source = "https://huggingface.co/Qwen/Qwen3.6-35B-A3B"
+18
View File
@@ -0,0 +1,18 @@
name = "Qwen3.6 Flash"
family = "qwen3.6"
release_date = "2026-04-27"
last_updated = "2026-04-27"
attachment = true
reasoning = true
temperature = true
tool_call = true
structured_output = true
open_weights = false
[limit]
context = 1_000_000
output = 65_536
[modalities]
input = ["text", "image", "video"]
output = ["text"]
+18
View File
@@ -0,0 +1,18 @@
name = "Qwen3.6 Max Preview"
family = "qwen"
release_date = "2026-04-20"
last_updated = "2026-04-20"
attachment = false
reasoning = true
temperature = true
tool_call = true
knowledge = "2025-04"
open_weights = false
[limit]
context = 262_144
output = 65_536
[modalities]
input = ["text"]
output = ["text"]
+18
View File
@@ -0,0 +1,18 @@
name = "Qwen3.6 Plus"
family = "qwen"
release_date = "2026-04-02"
last_updated = "2026-04-02"
attachment = false
reasoning = true
temperature = true
tool_call = true
knowledge = "2025-04"
open_weights = false
[limit]
context = 1_000_000
output = 65_536
[modalities]
input = ["text", "image", "video"]
output = ["text"]
+17
View File
@@ -0,0 +1,17 @@
name = "Qwen3.7 Max"
family = "qwen"
release_date = "2026-05-21"
last_updated = "2026-05-21"
attachment = false
reasoning = true
temperature = true
tool_call = true
open_weights = false
[limit]
context = 1_000_000
output = 65_536
[modalities]
input = ["text"]
output = ["text"]
+18
View File
@@ -0,0 +1,18 @@
name = "Qwen3.7 Plus"
family = "qwen"
release_date = "2026-06-02"
last_updated = "2026-06-02"
attachment = false
reasoning = true
temperature = true
tool_call = true
knowledge = "2025-04"
open_weights = false
[limit]
context = 131_072
output = 16_384
[modalities]
input = ["text"]
output = ["text"]
+18
View File
@@ -0,0 +1,18 @@
name = "QwQ Plus"
family = "qwen"
release_date = "2025-03-05"
last_updated = "2025-03-05"
attachment = false
reasoning = true
temperature = true
tool_call = true
knowledge = "2024-04"
open_weights = false
[limit]
context = 131_072
output = 8_192
[modalities]
input = ["text"]
output = ["text"]
@@ -0,0 +1,25 @@
name = "Claude Haiku 3.5"
family = "claude-haiku"
release_date = "2024-10-22"
last_updated = "2024-10-22"
attachment = true
reasoning = false
temperature = true
tool_call = true
knowledge = "2024-07-31"
open_weights = false
[limit]
context = 200_000
output = 8_192
[modalities]
input = ["text", "image", "pdf"]
output = ["text"]
[[benchmarks]]
name = "Aider Polyglot"
score = 28.0
metric = "percent correct"
source = "https://aider.chat/docs/leaderboards/"
date = "2024-12-21"
@@ -0,0 +1,25 @@
name = "Claude Sonnet 3.5 v2"
family = "claude-sonnet"
release_date = "2024-10-22"
last_updated = "2024-10-22"
attachment = true
reasoning = false
temperature = true
tool_call = true
knowledge = "2024-04-30"
open_weights = false
[limit]
context = 200_000
output = 8_192
[modalities]
input = ["text", "image", "pdf"]
output = ["text"]
[[benchmarks]]
name = "Aider Polyglot"
score = 51.6
metric = "percent correct"
source = "https://aider.chat/docs/leaderboards/"
date = "2025-01-17"
@@ -0,0 +1,25 @@
name = "Claude Sonnet 3.7"
family = "claude-sonnet"
release_date = "2025-02-19"
last_updated = "2025-02-19"
attachment = true
reasoning = true
temperature = true
tool_call = true
knowledge = "2024-10-31"
open_weights = false
[limit]
context = 200_000
output = 64_000
[modalities]
input = ["text", "image", "pdf"]
output = ["text"]
[[benchmarks]]
name = "Aider Polyglot"
score = 64.9
metric = "percent correct"
source = "https://aider.chat/docs/leaderboards/"
date = "2025-02-24"
@@ -1,19 +1,16 @@
name = "Claude 3.7 Sonnet Thinking (8K)"
release_date = "2025-02-24"
last_updated = "2025-02-24"
name = "Claude Haiku 4.5"
family = "claude-haiku"
release_date = "2025-10-15"
last_updated = "2025-10-15"
attachment = true
reasoning = true
temperature = true
tool_call = true
structured_output = true
knowledge = "2025-02-28"
open_weights = false
[cost]
input = 2.992
output = 14.994
[limit]
context = 200_000
input = 200_000
output = 64_000
[modalities]
+25
View File
@@ -0,0 +1,25 @@
name = "Claude Haiku 4.5 (latest)"
family = "claude-haiku"
release_date = "2025-10-15"
last_updated = "2025-10-15"
attachment = true
reasoning = true
temperature = true
tool_call = true
knowledge = "2025-02-28"
open_weights = false
[limit]
context = 200_000
output = 64_000
[modalities]
input = ["text", "image", "pdf"]
output = ["text"]
[[benchmarks]]
name = "SWE-Bench Pro"
score = 39.45
metric = "resolve rate"
dataset = "public"
source = "https://labs.scale.com/leaderboard/swe_bench_pro_public"
+25
View File
@@ -0,0 +1,25 @@
name = "Claude Opus 4 (latest)"
family = "claude-opus"
release_date = "2025-05-22"
last_updated = "2025-05-22"
attachment = true
reasoning = true
temperature = true
tool_call = true
knowledge = "2025-03-31"
open_weights = false
[limit]
context = 200_000
output = 32_000
[modalities]
input = ["text", "image", "pdf"]
output = ["text"]
[[benchmarks]]
name = "Aider Polyglot"
score = 72.0
metric = "percent correct"
source = "https://aider.chat/docs/leaderboards/"
date = "2025-05-25"
@@ -1,23 +1,18 @@
name = "Claude Opus 4.1"
family = "claude-opus"
status = "deprecated"
release_date = "2025-08-05"
last_updated = "2025-08-05"
attachment = true
reasoning = true
temperature = true
tool_call = false
tool_call = true
knowledge = "2025-03-31"
open_weights = false
[cost]
input = 0
output = 0
[limit]
context = 80_000
output = 16_000
context = 200_000
output = 32_000
[modalities]
input = ["text", "image"]
input = ["text", "image", "pdf"]
output = ["text"]
+18
View File
@@ -0,0 +1,18 @@
name = "Claude Opus 4.1 (latest)"
family = "claude-opus"
release_date = "2025-08-05"
last_updated = "2025-08-05"
attachment = true
reasoning = true
temperature = true
tool_call = true
knowledge = "2025-03-31"
open_weights = false
[limit]
context = 200_000
output = 32_000
[modalities]
input = ["text", "image", "pdf"]
output = ["text"]
@@ -0,0 +1,25 @@
name = "Claude Opus 4"
family = "claude-opus"
release_date = "2025-05-22"
last_updated = "2025-05-22"
attachment = true
reasoning = true
temperature = true
tool_call = true
knowledge = "2025-03-31"
open_weights = false
[limit]
context = 200_000
output = 32_000
[modalities]
input = ["text", "image", "pdf"]
output = ["text"]
[[benchmarks]]
name = "Aider Polyglot"
score = 72.0
metric = "percent correct"
source = "https://aider.chat/docs/leaderboards/"
date = "2025-05-25"
@@ -0,0 +1,25 @@
name = "Claude Opus 4.5"
family = "claude-opus"
release_date = "2025-11-01"
last_updated = "2025-11-01"
attachment = true
reasoning = true
temperature = true
tool_call = true
knowledge = "2025-03-31"
open_weights = false
[limit]
context = 200_000
output = 64_000
[modalities]
input = ["text", "image", "pdf"]
output = ["text"]
[[benchmarks]]
name = "SWE-Bench Pro"
score = 45.89
metric = "resolve rate"
dataset = "public"
source = "https://labs.scale.com/leaderboard/swe_bench_pro_public"
@@ -1,19 +1,16 @@
name = "Claude 3.7 Sonnet Thinking (1K)"
release_date = "2025-02-24"
last_updated = "2025-02-24"
name = "Claude Opus 4.5 (latest)"
family = "claude-opus"
release_date = "2025-11-24"
last_updated = "2025-11-24"
attachment = true
reasoning = true
temperature = true
tool_call = true
structured_output = true
knowledge = "2025-03-31"
open_weights = false
[cost]
input = 2.992
output = 14.994
[limit]
context = 200_000
input = 200_000
output = 64_000
[modalities]
+94
View File
@@ -0,0 +1,94 @@
name = "Claude Opus 4.6"
family = "claude-opus"
release_date = "2026-02-05"
last_updated = "2026-03-13"
attachment = true
reasoning = true
temperature = true
tool_call = true
knowledge = "2025-05-31"
open_weights = false
[limit]
context = 1_000_000
output = 128_000
[modalities]
input = ["text", "image", "pdf"]
output = ["text"]
[[benchmarks]]
name = "SWE-Bench Pro"
score = 51.9
metric = "resolve rate"
dataset = "public"
source = "https://labs.scale.com/leaderboard/swe_bench_pro_public"
[[benchmarks]]
name = "SWE-Atlas Codebase QnA"
score = 33.3
metric = "score"
harness = "Claude Code"
source = "https://labs.scale.com/leaderboard/sweatlas-qna"
[[benchmarks]]
name = "SWE-Atlas Codebase QnA"
score = 30
metric = "score"
harness = "Mini-SWE-Agent"
source = "https://labs.scale.com/leaderboard/sweatlas-qna"
[[benchmarks]]
name = "SWE-Atlas Refactoring"
score = 35.58
metric = "score"
harness = "Claude Code"
source = "https://labs.scale.com/leaderboard/sweatlas-refactoring"
[[benchmarks]]
name = "SWE-Atlas Test Writing"
score = 36.67
metric = "score"
harness = "Claude Code"
source = "https://labs.scale.com/leaderboard/sweatlas-tw"
[[benchmarks]]
name = "SWE-Atlas Test Writing"
score = 36.08
metric = "score"
harness = "Mini-SWE-Agent"
source = "https://labs.scale.com/leaderboard/sweatlas-tw"
[[benchmarks]]
name = "Artificial Analysis Coding Agent Index"
score = 51.3
metric = "average pass@1"
harness = "Claude Code"
variant = "medium"
source = "https://artificialanalysis.ai/agents/coding-agents"
[[benchmarks]]
name = "SWE-Atlas Codebase QnA"
score = 71.9
metric = "pass@1"
harness = "Claude Code"
variant = "medium"
source = "https://artificialanalysis.ai/agents/coding-agents"
[[benchmarks]]
name = "SWE-Bench Pro"
score = 11.8
metric = "pass@1"
harness = "Claude Code"
variant = "medium"
dataset = "hard-aa"
source = "https://artificialanalysis.ai/agents/coding-agents"
[[benchmarks]]
name = "Terminal-Bench"
score = 70.2
metric = "pass@1"
harness = "Claude Code"
variant = "medium"
version = "2.1"
source = "https://artificialanalysis.ai/agents/coding-agents"
+143
View File
@@ -0,0 +1,143 @@
name = "Claude Opus 4.7"
family = "claude-opus"
release_date = "2026-04-16"
last_updated = "2026-04-16"
attachment = true
reasoning = true
temperature = false
tool_call = true
knowledge = "2026-01-31"
open_weights = false
[limit]
context = 1_000_000
output = 128_000
[modalities]
input = ["text", "image", "pdf"]
output = ["text"]
[[benchmarks]]
name = "SWE-Bench Pro"
score = 64.3
metric = "resolve rate"
source = "https://www.anthropic.com/news/claude-opus-4-8"
date = "2026-05-28"
[[benchmarks]]
name = "Terminal-Bench"
score = 66.1
metric = "success rate"
harness = "Terminus-2"
version = "2.1"
source = "https://www.anthropic.com/news/claude-opus-4-8"
date = "2026-05-28"
[[benchmarks]]
name = "SWE-Atlas Refactoring"
score = 48.57
metric = "score"
harness = "Claude Code"
source = "https://labs.scale.com/leaderboard/sweatlas-refactoring"
[[benchmarks]]
name = "Artificial Analysis Coding Agent Index"
score = 66.6
metric = "average pass@1"
harness = "Claude Code"
variant = "max"
source = "https://artificialanalysis.ai/agents/coding-agents"
[[benchmarks]]
name = "SWE-Atlas Codebase QnA"
score = 81
metric = "pass@1"
harness = "Claude Code"
variant = "max"
source = "https://artificialanalysis.ai/agents/coding-agents"
[[benchmarks]]
name = "SWE-Bench Pro"
score = 44.9
metric = "pass@1"
harness = "Claude Code"
variant = "max"
dataset = "hard-aa"
source = "https://artificialanalysis.ai/agents/coding-agents"
[[benchmarks]]
name = "Terminal-Bench"
score = 73.8
metric = "pass@1"
harness = "Claude Code"
variant = "max"
version = "2.1"
source = "https://artificialanalysis.ai/agents/coding-agents"
[[benchmarks]]
name = "Artificial Analysis Coding Agent Index"
score = 61.2
metric = "average pass@1"
harness = "Cursor CLI"
variant = "medium"
source = "https://artificialanalysis.ai/agents/coding-agents"
[[benchmarks]]
name = "SWE-Atlas Codebase QnA"
score = 78.4
metric = "pass@1"
harness = "Cursor CLI"
variant = "medium"
source = "https://artificialanalysis.ai/agents/coding-agents"
[[benchmarks]]
name = "SWE-Bench Pro"
score = 34.4
metric = "pass@1"
harness = "Cursor CLI"
variant = "medium"
dataset = "hard-aa"
source = "https://artificialanalysis.ai/agents/coding-agents"
[[benchmarks]]
name = "Terminal-Bench"
score = 70.6
metric = "pass@1"
harness = "Cursor CLI"
variant = "medium"
version = "2.1"
source = "https://artificialanalysis.ai/agents/coding-agents"
[[benchmarks]]
name = "Artificial Analysis Coding Agent Index"
score = 59.9
metric = "average pass@1"
harness = "Claude Code"
variant = "medium"
source = "https://artificialanalysis.ai/agents/coding-agents"
[[benchmarks]]
name = "SWE-Atlas Codebase QnA"
score = 71.7
metric = "pass@1"
harness = "Claude Code"
variant = "medium"
source = "https://artificialanalysis.ai/agents/coding-agents"
[[benchmarks]]
name = "SWE-Bench Pro"
score = 36.4
metric = "pass@1"
harness = "Claude Code"
variant = "medium"
dataset = "hard-aa"
source = "https://artificialanalysis.ai/agents/coding-agents"
[[benchmarks]]
name = "Terminal-Bench"
score = 71.4
metric = "pass@1"
harness = "Claude Code"
variant = "medium"
version = "2.1"
source = "https://artificialanalysis.ai/agents/coding-agents"
+33
View File
@@ -0,0 +1,33 @@
name = "Claude Opus 4.8"
family = "claude-opus"
release_date = "2026-05-28"
last_updated = "2026-05-28"
attachment = true
reasoning = true
temperature = false
tool_call = true
open_weights = false
[limit]
context = 1_000_000
output = 128_000
[modalities]
input = ["text", "image", "pdf"]
output = ["text"]
[[benchmarks]]
name = "SWE-Bench Pro"
score = 69.2
metric = "resolve rate"
source = "https://www.anthropic.com/news/claude-opus-4-8"
date = "2026-05-28"
[[benchmarks]]
name = "Terminal-Bench"
score = 74.6
metric = "success rate"
harness = "Terminus-2"
version = "2.1"
source = "https://www.anthropic.com/news/claude-opus-4-8"
date = "2026-05-28"
+32
View File
@@ -0,0 +1,32 @@
name = "Claude Sonnet 4 (latest)"
family = "claude-sonnet"
release_date = "2025-05-22"
last_updated = "2025-05-22"
attachment = true
reasoning = true
temperature = true
tool_call = true
knowledge = "2025-03-31"
open_weights = false
[limit]
context = 200_000
output = 64_000
[modalities]
input = ["text", "image", "pdf"]
output = ["text"]
[[benchmarks]]
name = "Aider Polyglot"
score = 61.3
metric = "percent correct"
source = "https://aider.chat/docs/leaderboards/"
date = "2025-05-24"
[[benchmarks]]
name = "SWE-Bench Pro"
score = 42.7
metric = "resolve rate"
dataset = "public"
source = "https://labs.scale.com/leaderboard/swe_bench_pro_public"
@@ -0,0 +1,25 @@
name = "Claude Sonnet 4"
family = "claude-sonnet"
release_date = "2025-05-22"
last_updated = "2025-05-22"
attachment = true
reasoning = true
temperature = true
tool_call = true
knowledge = "2025-03-31"
open_weights = false
[limit]
context = 200_000
output = 64_000
[modalities]
input = ["text", "image", "pdf"]
output = ["text"]
[[benchmarks]]
name = "Aider Polyglot"
score = 61.3
metric = "percent correct"
source = "https://aider.chat/docs/leaderboards/"
date = "2025-05-24"
@@ -1,19 +1,16 @@
name = "Claude 3.7 Sonnet Thinking (32K)"
release_date = "2025-07-15"
last_updated = "2025-07-15"
name = "Claude Sonnet 4.5"
family = "claude-sonnet"
release_date = "2025-09-29"
last_updated = "2025-09-29"
attachment = true
reasoning = true
temperature = true
tool_call = true
structured_output = true
knowledge = "2025-07-31"
open_weights = false
[cost]
input = 2.992
output = 14.994
[limit]
context = 200_000
input = 200_000
output = 64_000
[modalities]
+25
View File
@@ -0,0 +1,25 @@
name = "Claude Sonnet 4.5 (latest)"
family = "claude-sonnet"
release_date = "2025-09-29"
last_updated = "2025-09-29"
attachment = true
reasoning = true
temperature = true
tool_call = true
knowledge = "2025-07-31"
open_weights = false
[limit]
context = 200_000
output = 64_000
[modalities]
input = ["text", "image", "pdf"]
output = ["text"]
[[benchmarks]]
name = "SWE-Bench Pro"
score = 43.6
metric = "resolve rate"
dataset = "public"
source = "https://labs.scale.com/leaderboard/swe_bench_pro_public"
+73
View File
@@ -0,0 +1,73 @@
name = "Claude Sonnet 4.6"
family = "claude-sonnet"
release_date = "2026-02-17"
last_updated = "2026-03-13"
attachment = true
reasoning = true
temperature = true
tool_call = true
knowledge = "2025-08-31"
open_weights = false
[limit]
context = 1_000_000
output = 64_000
[modalities]
input = ["text", "image", "pdf"]
output = ["text"]
[[benchmarks]]
name = "SWE-Atlas Codebase QnA"
score = 31.2
metric = "score"
harness = "Claude Code"
source = "https://labs.scale.com/leaderboard/sweatlas-qna"
[[benchmarks]]
name = "SWE-Atlas Refactoring"
score = 32.21
metric = "score"
harness = "Claude Code"
source = "https://labs.scale.com/leaderboard/sweatlas-refactoring"
[[benchmarks]]
name = "SWE-Atlas Test Writing"
score = 31.76
metric = "score"
harness = "Claude Code"
source = "https://labs.scale.com/leaderboard/sweatlas-tw"
[[benchmarks]]
name = "Artificial Analysis Coding Agent Index"
score = 49.4
metric = "average pass@1"
harness = "Claude Code"
variant = "medium"
source = "https://artificialanalysis.ai/agents/coding-agents"
[[benchmarks]]
name = "SWE-Atlas Codebase QnA"
score = 70.3
metric = "pass@1"
harness = "Claude Code"
variant = "medium"
source = "https://artificialanalysis.ai/agents/coding-agents"
[[benchmarks]]
name = "SWE-Bench Pro"
score = 14.9
metric = "pass@1"
harness = "Claude Code"
variant = "medium"
dataset = "hard-aa"
source = "https://artificialanalysis.ai/agents/coding-agents"
[[benchmarks]]
name = "Terminal-Bench"
score = 63.1
metric = "pass@1"
harness = "Claude Code"
variant = "medium"
version = "2.1"
source = "https://artificialanalysis.ai/agents/coding-agents"
+29
View File
@@ -0,0 +1,29 @@
name = "Command A"
family = "command-a"
release_date = "2025-03-13"
last_updated = "2025-03-13"
attachment = false
reasoning = false
temperature = true
tool_call = true
knowledge = "2024-06-01"
open_weights = true
[limit]
context = 256_000
output = 8_000
[modalities]
input = ["text"]
output = ["text"]
[[weights]]
label = "Hugging Face"
url = "https://huggingface.co/CohereLabs/c4ai-command-a-03-2025"
[[benchmarks]]
name = "Aider Polyglot"
score = 12.0
metric = "percent correct"
source = "https://aider.chat/docs/leaderboards/"
date = "2025-03-14"
+22
View File
@@ -0,0 +1,22 @@
name = "Command R"
family = "command-r"
release_date = "2024-08-30"
last_updated = "2024-08-30"
attachment = false
reasoning = false
temperature = true
tool_call = true
knowledge = "2024-06-01"
open_weights = true
[limit]
context = 128_000
output = 4_000
[modalities]
input = ["text"]
output = ["text"]
[[weights]]
label = "Hugging Face"
url = "https://huggingface.co/CohereLabs/c4ai-command-r-08-2024"
+22
View File
@@ -0,0 +1,22 @@
name = "Command R+"
family = "command-r"
release_date = "2024-08-30"
last_updated = "2024-08-30"
attachment = false
reasoning = false
temperature = true
tool_call = true
knowledge = "2024-06-01"
open_weights = true
[limit]
context = 128_000
output = 4_000
[modalities]
input = ["text"]
output = ["text"]
[[weights]]
label = "Hugging Face"
url = "https://huggingface.co/CohereLabs/c4ai-command-r-plus-08-2024"
+22
View File
@@ -0,0 +1,22 @@
name = "Command R7B"
family = "command-r"
release_date = "2024-02-27"
last_updated = "2024-02-27"
attachment = false
reasoning = false
temperature = true
tool_call = true
knowledge = "2024-06-01"
open_weights = true
[limit]
context = 128_000
output = 4_000
[modalities]
input = ["text"]
output = ["text"]
[[weights]]
label = "Hugging Face"
url = "https://huggingface.co/CohereLabs/c4ai-command-r7b-12-2024"
+29
View File
@@ -0,0 +1,29 @@
name = "DeepSeek Chat"
family = "deepseek"
release_date = "2025-12-01"
last_updated = "2026-02-28"
attachment = true
reasoning = false
temperature = true
tool_call = true
knowledge = "2025-09"
open_weights = true
[limit]
context = 1_000_000
output = 384_000
[modalities]
input = ["text"]
output = ["text"]
[[weights]]
label = "Hugging Face"
url = "https://huggingface.co/deepseek-ai/DeepSeek-V3.2"
[[benchmarks]]
name = "Aider Polyglot"
score = 70.2
metric = "percent correct"
source = "https://aider.chat/docs/leaderboards/"
date = "2025-10-03"
+50
View File
@@ -0,0 +1,50 @@
name = "DeepSeek-R1"
family = "deepseek-thinking"
release_date = "2025-01-20"
last_updated = "2025-05-29"
attachment = false
reasoning = true
temperature = true
tool_call = true
knowledge = "2024-07"
open_weights = true
[limit]
context = 128_000
output = 32_768
[modalities]
input = ["text"]
output = ["text"]
[[weights]]
label = "Hugging Face"
url = "https://huggingface.co/deepseek-ai/DeepSeek-R1"
[[benchmarks]]
name = "Aider Polyglot"
score = 56.9
metric = "percent correct"
source = "https://aider.chat/docs/leaderboards/"
date = "2025-01-20"
[[benchmarks]]
name = "Artificial Analysis Coding Index"
score = 15.9
metric = "index"
source = "https://openrouter.ai/deepseek/deepseek-r1/benchmarks"
date = "2026-03-11"
[[benchmarks]]
name = "SciCode"
score = 35.7
metric = "percent correct"
source = "https://openrouter.ai/deepseek/deepseek-r1/benchmarks"
date = "2026-03-11"
[[benchmarks]]
name = "Terminal-Bench Hard"
score = 6.1
metric = "success rate"
source = "https://openrouter.ai/deepseek/deepseek-r1/benchmarks"
date = "2026-03-11"
+29
View File
@@ -0,0 +1,29 @@
name = "DeepSeek Reasoner"
family = "deepseek-thinking"
release_date = "2025-12-01"
last_updated = "2026-02-28"
attachment = true
reasoning = true
temperature = true
tool_call = true
knowledge = "2025-09"
open_weights = true
[limit]
context = 1_000_000
output = 384_000
[modalities]
input = ["text"]
output = ["text"]
[[weights]]
label = "Hugging Face"
url = "https://huggingface.co/deepseek-ai/DeepSeek-V3.2"
[[benchmarks]]
name = "Aider Polyglot"
score = 74.2
metric = "percent correct"
source = "https://aider.chat/docs/leaderboards/"
date = "2025-10-03"
+29
View File
@@ -0,0 +1,29 @@
name = "DeepSeek V4 Flash"
family = "deepseek-flash"
release_date = "2026-04-24"
last_updated = "2026-04-24"
attachment = false
reasoning = true
temperature = true
tool_call = true
structured_output = true
knowledge = "2025-05"
open_weights = true
[limit]
context = 1_000_000
output = 384_000
[modalities]
input = ["text"]
output = ["text"]
[[weights]]
label = "Hugging Face"
url = "https://huggingface.co/deepseek-ai/DeepSeek-V4-Flash"
[[benchmarks]]
name = "SWE-Bench Verified"
score = 79
metric = "resolved"
source = "https://huggingface.co/deepseek-ai/DeepSeek-V4-Flash"
+63
View File
@@ -0,0 +1,63 @@
name = "DeepSeek V4 Pro"
family = "deepseek-thinking"
release_date = "2026-04-24"
last_updated = "2026-04-24"
attachment = false
reasoning = true
temperature = true
tool_call = true
structured_output = true
knowledge = "2025-05"
open_weights = true
[limit]
context = 1_000_000
output = 384_000
[modalities]
input = ["text"]
output = ["text"]
[[weights]]
label = "Hugging Face"
url = "https://huggingface.co/deepseek-ai/DeepSeek-V4-Pro"
[[benchmarks]]
name = "SWE-Bench Verified"
score = 80.6
metric = "resolved"
source = "https://huggingface.co/deepseek-ai/DeepSeek-V4-Pro"
[[benchmarks]]
name = "Artificial Analysis Coding Agent Index"
score = 50.1
metric = "average pass@1"
harness = "Claude Code"
variant = "high"
source = "https://artificialanalysis.ai/agents/coding-agents"
[[benchmarks]]
name = "SWE-Atlas Codebase QnA"
score = 67.8
metric = "pass@1"
harness = "Claude Code"
variant = "high"
source = "https://artificialanalysis.ai/agents/coding-agents"
[[benchmarks]]
name = "SWE-Bench Pro"
score = 18
metric = "pass@1"
harness = "Claude Code"
variant = "high"
dataset = "hard-aa"
source = "https://artificialanalysis.ai/agents/coding-agents"
[[benchmarks]]
name = "Terminal-Bench"
score = 64.7
metric = "pass@1"
harness = "Claude Code"
variant = "high"
version = "2.1"
source = "https://artificialanalysis.ai/agents/coding-agents"
+19
View File
@@ -0,0 +1,19 @@
name = "Gemini 2.0 Flash-Lite"
family = "gemini-flash-lite"
release_date = "2024-12-11"
last_updated = "2024-12-11"
attachment = true
reasoning = false
temperature = true
tool_call = true
structured_output = true
knowledge = "2024-06"
open_weights = false
[limit]
context = 1_048_576
output = 8_192
[modalities]
input = ["text", "image", "audio", "video", "pdf"]
output = ["text"]
+19
View File
@@ -0,0 +1,19 @@
name = "Gemini 2.0 Flash"
family = "gemini-flash"
release_date = "2024-12-11"
last_updated = "2024-12-11"
attachment = true
reasoning = false
temperature = true
tool_call = true
structured_output = true
knowledge = "2024-06"
open_weights = false
[limit]
context = 1_048_576
output = 8_192
[modalities]
input = ["text", "image", "audio", "video", "pdf"]
output = ["text"]
+18
View File
@@ -0,0 +1,18 @@
name = "Nano Banana"
family = "gemini-flash"
release_date = "2025-08-26"
last_updated = "2025-08-26"
attachment = true
reasoning = true
temperature = true
tool_call = false
knowledge = "2025-06"
open_weights = false
[limit]
context = 32_768
output = 32_768
[modalities]
input = ["text", "image"]
output = ["text", "image"]
+40
View File
@@ -0,0 +1,40 @@
name = "Gemini 2.5 Flash-Lite"
family = "gemini-flash-lite"
release_date = "2025-06-17"
last_updated = "2025-06-17"
attachment = true
reasoning = true
temperature = true
tool_call = true
structured_output = true
knowledge = "2025-01"
open_weights = false
[limit]
context = 1_048_576
output = 65_536
[modalities]
input = ["text", "image", "audio", "video", "pdf"]
output = ["text"]
[[benchmarks]]
name = "Artificial Analysis Coding Index"
score = 9.5
metric = "index"
source = "https://openrouter.ai/google/gemini-2.5-flash-lite/benchmarks"
date = "2026-03-11"
[[benchmarks]]
name = "SciCode"
score = 19.3
metric = "percent correct"
source = "https://openrouter.ai/google/gemini-2.5-flash-lite/benchmarks"
date = "2026-03-11"
[[benchmarks]]
name = "Terminal-Bench Hard"
score = 4.5
metric = "success rate"
source = "https://openrouter.ai/google/gemini-2.5-flash-lite/benchmarks"
date = "2026-03-11"
+47
View File
@@ -0,0 +1,47 @@
name = "Gemini 2.5 Flash"
family = "gemini-flash"
release_date = "2025-03-20"
last_updated = "2025-06-05"
attachment = true
reasoning = true
temperature = true
tool_call = true
structured_output = true
knowledge = "2025-01"
open_weights = false
[limit]
context = 1_048_576
output = 65_536
[modalities]
input = ["text", "image", "audio", "video", "pdf"]
output = ["text"]
[[benchmarks]]
name = "Aider Polyglot"
score = 55.1
metric = "percent correct"
source = "https://aider.chat/docs/leaderboards/"
date = "2025-05-25"
[[benchmarks]]
name = "Artificial Analysis Coding Index"
score = 22.2
metric = "index"
source = "https://openrouter.ai/google/gemini-2.5-flash/benchmarks"
date = "2026-06-02"
[[benchmarks]]
name = "SciCode"
score = 39.4
metric = "percent correct"
source = "https://openrouter.ai/google/gemini-2.5-flash/benchmarks"
date = "2026-06-02"
[[benchmarks]]
name = "Terminal-Bench Hard"
score = 13.6
metric = "success rate"
source = "https://openrouter.ai/google/gemini-2.5-flash/benchmarks"
date = "2026-06-02"
+47
View File
@@ -0,0 +1,47 @@
name = "Gemini 2.5 Pro"
family = "gemini-pro"
release_date = "2025-03-20"
last_updated = "2025-06-05"
attachment = true
reasoning = true
temperature = true
tool_call = true
structured_output = true
knowledge = "2025-01"
open_weights = false
[limit]
context = 1_048_576
output = 65_536
[modalities]
input = ["text", "image", "audio", "video", "pdf"]
output = ["text"]
[[benchmarks]]
name = "Aider Polyglot"
score = 83.1
metric = "percent correct"
source = "https://aider.chat/docs/leaderboards/"
date = "2025-06-06"
[[benchmarks]]
name = "Artificial Analysis Coding Index"
score = 32
metric = "index"
source = "https://openrouter.ai/google/gemini-2.5-pro/benchmarks"
date = "2026-06-02"
[[benchmarks]]
name = "SciCode"
score = 42.8
metric = "percent correct"
source = "https://openrouter.ai/google/gemini-2.5-pro/benchmarks"
date = "2026-06-02"
[[benchmarks]]
name = "Terminal-Bench Hard"
score = 26.5
metric = "success rate"
source = "https://openrouter.ai/google/gemini-2.5-pro/benchmarks"
date = "2026-06-02"
+47
View File
@@ -0,0 +1,47 @@
name = "Gemini 3 Flash Preview"
family = "gemini-flash"
release_date = "2025-12-17"
last_updated = "2025-12-17"
attachment = true
reasoning = true
temperature = true
tool_call = true
structured_output = true
knowledge = "2025-01"
open_weights = false
[limit]
context = 1_048_576
output = 65_536
[modalities]
input = ["text", "image", "video", "audio", "pdf"]
output = ["text"]
[[benchmarks]]
name = "SWE-Bench Pro"
score = 34.63
metric = "resolve rate"
dataset = "public"
source = "https://labs.scale.com/leaderboard/swe_bench_pro_public"
[[benchmarks]]
name = "SWE-Atlas Codebase QnA"
score = 8.2
metric = "score"
harness = "Mini-SWE-Agent"
source = "https://labs.scale.com/leaderboard/sweatlas-qna"
[[benchmarks]]
name = "SWE-Atlas Refactoring"
score = 10
metric = "score"
harness = "Mini-SWE-Agent"
source = "https://labs.scale.com/leaderboard/sweatlas-refactoring"
[[benchmarks]]
name = "SWE-Atlas Test Writing"
score = 30.3
metric = "score"
harness = "Mini-SWE-Agent"
source = "https://labs.scale.com/leaderboard/sweatlas-tw"
@@ -1,6 +1,5 @@
name = "Gemini 3 Pro Preview"
family = "gemini-pro"
status = "deprecated"
release_date = "2025-11-18"
last_updated = "2025-11-18"
attachment = true
@@ -11,15 +10,17 @@ structured_output = true
knowledge = "2025-01"
open_weights = false
[cost]
input = 0
output = 0
[limit]
context = 128_000
output = 64_000
input = 128_000
context = 1_048_576
output = 65_536
[modalities]
input = ["text", "image", "audio", "video"]
input = ["text", "image", "video", "audio", "pdf"]
output = ["text"]
[[benchmarks]]
name = "SWE-Bench Pro"
score = 43.3
metric = "resolve rate"
dataset = "public"
source = "https://labs.scale.com/leaderboard/swe_bench_pro_public"
@@ -0,0 +1,18 @@
name = "Nano Banana 2"
family = "gemini-flash"
release_date = "2026-02-26"
last_updated = "2026-02-26"
attachment = true
reasoning = true
temperature = true
tool_call = false
knowledge = "2025-01"
open_weights = false
[limit]
context = 65_536
output = 65_536
[modalities]
input = ["text", "image", "pdf"]
output = ["text", "image"]
@@ -0,0 +1,19 @@
name = "Gemini 3.1 Flash Lite Preview"
family = "gemini-flash-lite"
release_date = "2026-03-03"
last_updated = "2026-03-03"
attachment = true
reasoning = true
temperature = true
tool_call = true
structured_output = true
knowledge = "2025-01"
open_weights = false
[limit]
context = 1_048_576
output = 65_536
[modalities]
input = ["text", "image", "video", "audio", "pdf"]
output = ["text"]
+19
View File
@@ -0,0 +1,19 @@
name = "Gemini 3.1 Flash Lite"
family = "gemini-flash-lite"
release_date = "2026-05-07"
last_updated = "2026-05-07"
attachment = true
reasoning = true
temperature = true
tool_call = true
structured_output = true
knowledge = "2025-01"
open_weights = false
[limit]
context = 1_048_576
output = 65_536
[modalities]
input = ["text", "image", "video", "audio", "pdf"]
output = ["text"]
@@ -0,0 +1,19 @@
name = "Gemini 3.1 Pro Preview Custom Tools"
family = "gemini-pro"
release_date = "2026-02-19"
last_updated = "2026-02-19"
attachment = true
reasoning = true
temperature = true
tool_call = true
structured_output = true
knowledge = "2025-01"
open_weights = false
[limit]
context = 1_048_576
output = 65_536
[modalities]
input = ["text", "image", "video", "audio", "pdf"]
output = ["text"]
+97
View File
@@ -0,0 +1,97 @@
name = "Gemini 3.1 Pro Preview"
family = "gemini-pro"
release_date = "2026-02-19"
last_updated = "2026-02-19"
attachment = true
reasoning = true
temperature = true
tool_call = true
structured_output = true
knowledge = "2025-01"
open_weights = false
[limit]
context = 1_048_576
output = 65_536
[modalities]
input = ["text", "image", "video", "audio", "pdf"]
output = ["text"]
[[benchmarks]]
name = "SWE-Bench Pro"
score = 54.2
metric = "resolve rate"
source = "https://www.anthropic.com/news/claude-opus-4-8"
date = "2026-05-28"
[[benchmarks]]
name = "Terminal-Bench"
score = 70.3
metric = "success rate"
harness = "Terminus-2"
version = "2.1"
source = "https://www.anthropic.com/news/claude-opus-4-8"
date = "2026-05-28"
[[benchmarks]]
name = "SWE-Bench Pro"
score = 46.1
metric = "resolve rate"
dataset = "public"
source = "https://labs.scale.com/leaderboard/swe_bench_pro_public"
[[benchmarks]]
name = "SWE-Atlas Codebase QnA"
score = 13.5
metric = "score"
harness = "Mini-SWE-Agent"
source = "https://labs.scale.com/leaderboard/sweatlas-qna"
[[benchmarks]]
name = "SWE-Atlas Refactoring"
score = 33.81
metric = "score"
harness = "Gemini CLI"
source = "https://labs.scale.com/leaderboard/sweatlas-refactoring"
[[benchmarks]]
name = "SWE-Atlas Test Writing"
score = 29.84
metric = "score"
harness = "Mini-SWE-Agent"
source = "https://labs.scale.com/leaderboard/sweatlas-tw"
[[benchmarks]]
name = "Artificial Analysis Coding Agent Index"
score = 43
metric = "average pass@1"
harness = "Gemini CLI"
variant = "high"
source = "https://artificialanalysis.ai/agents/coding-agents"
[[benchmarks]]
name = "SWE-Atlas Codebase QnA"
score = 45.6
metric = "pass@1"
harness = "Gemini CLI"
variant = "high"
source = "https://artificialanalysis.ai/agents/coding-agents"
[[benchmarks]]
name = "SWE-Bench Pro"
score = 15.1
metric = "pass@1"
harness = "Gemini CLI"
variant = "high"
dataset = "hard-aa"
source = "https://artificialanalysis.ai/agents/coding-agents"
[[benchmarks]]
name = "Terminal-Bench"
score = 68.3
metric = "pass@1"
harness = "Gemini CLI"
variant = "high"
version = "2.1"
source = "https://artificialanalysis.ai/agents/coding-agents"
+19
View File
@@ -0,0 +1,19 @@
name = "Gemini 3.5 Flash"
family = "gemini-flash"
release_date = "2026-05-19"
last_updated = "2026-05-19"
attachment = true
reasoning = true
temperature = true
tool_call = true
structured_output = true
knowledge = "2025-01"
open_weights = false
[limit]
context = 1_048_576
output = 65_536
[modalities]
input = ["text", "image", "video", "audio", "pdf"]
output = ["text"]
+18
View File
@@ -0,0 +1,18 @@
name = "Gemini Embedding 001"
family = "gemini"
release_date = "2025-05-20"
last_updated = "2025-05-20"
attachment = false
reasoning = false
temperature = false
tool_call = false
knowledge = "2025-05"
open_weights = false
[limit]
context = 2_048
output = 1
[modalities]
input = ["text"]
output = ["text"]
+19
View File
@@ -0,0 +1,19 @@
name = "Gemini Flash Latest"
family = "gemini-flash"
release_date = "2025-09-25"
last_updated = "2025-09-25"
attachment = true
reasoning = true
temperature = true
tool_call = true
structured_output = true
knowledge = "2025-01"
open_weights = false
[limit]
context = 1_048_576
output = 65_536
[modalities]
input = ["text", "image", "audio", "video", "pdf"]
output = ["text"]
@@ -0,0 +1,19 @@
name = "Gemini Flash-Lite Latest"
family = "gemini-flash-lite"
release_date = "2025-09-25"
last_updated = "2025-09-25"
attachment = true
reasoning = true
temperature = true
tool_call = true
structured_output = true
knowledge = "2025-01"
open_weights = false
[limit]
context = 1_048_576
output = 65_536
[modalities]
input = ["text", "image", "audio", "video", "pdf"]
output = ["text"]
+22
View File
@@ -0,0 +1,22 @@
name = "Gemma 4 26B A4B IT"
family = "gemma"
release_date = "2026-04-02"
last_updated = "2026-04-02"
attachment = true
reasoning = true
temperature = true
tool_call = true
structured_output = true
open_weights = true
[limit]
context = 262_144
output = 32_768
[modalities]
input = ["text", "image"]
output = ["text"]
[[weights]]
label = "Hugging Face"
url = "https://huggingface.co/google/gemma-4-26B-A4B-it"
+22
View File
@@ -0,0 +1,22 @@
name = "Gemma 4 31B IT"
family = "gemma"
release_date = "2026-04-02"
last_updated = "2026-04-02"
attachment = true
reasoning = true
temperature = true
tool_call = true
structured_output = true
open_weights = true
[limit]
context = 262_144
output = 32_768
[modalities]
input = ["text", "image"]
output = ["text"]
[[weights]]
label = "Hugging Face"
url = "https://huggingface.co/google/gemma-4-31B-it"
+43
View File
@@ -0,0 +1,43 @@
name = "Llama-3.3-70B-Instruct"
family = "llama"
release_date = "2024-12-06"
last_updated = "2024-12-06"
attachment = true
reasoning = false
temperature = true
tool_call = true
knowledge = "2023-12"
open_weights = true
[limit]
context = 128_000
output = 4_096
[modalities]
input = ["text"]
output = ["text"]
[[weights]]
label = "Hugging Face"
url = "https://huggingface.co/meta-llama/Llama-3.3-70B-Instruct"
[[benchmarks]]
name = "Artificial Analysis Coding Index"
score = 10.7
metric = "index"
source = "https://openrouter.ai/meta-llama/llama-3.3-70b-instruct/benchmarks"
date = "2026-03-11"
[[benchmarks]]
name = "SciCode"
score = 26
metric = "percent correct"
source = "https://openrouter.ai/meta-llama/llama-3.3-70b-instruct/benchmarks"
date = "2026-03-11"
[[benchmarks]]
name = "Terminal-Bench Hard"
score = 3
metric = "success rate"
source = "https://openrouter.ai/meta-llama/llama-3.3-70b-instruct/benchmarks"
date = "2026-03-11"
@@ -0,0 +1,36 @@
name = "Llama 4 Maverick 17B Instruct"
family = "llama"
release_date = "2025-04-05"
last_updated = "2025-04-05"
attachment = true
reasoning = false
temperature = true
tool_call = true
knowledge = "2024-08"
open_weights = true
[limit]
context = 1_000_000
output = 16_384
[modalities]
input = ["text", "image"]
output = ["text"]
[[weights]]
label = "Hugging Face"
url = "https://huggingface.co/meta-llama/Llama-4-Maverick-17B-128E-Instruct"
[[benchmarks]]
name = "Aider Polyglot"
score = 15.6
metric = "percent correct"
source = "https://aider.chat/docs/leaderboards/"
date = "2025-04-06"
[[benchmarks]]
name = "SWE-Bench Pro"
score = 5.24
metric = "resolve rate"
dataset = "public"
source = "https://labs.scale.com/leaderboard/swe_bench_pro_public"
@@ -0,0 +1,22 @@
name = "Llama 4 Scout 17B Instruct"
family = "llama"
release_date = "2025-04-05"
last_updated = "2025-04-05"
attachment = true
reasoning = false
temperature = true
tool_call = true
knowledge = "2024-08"
open_weights = true
[limit]
context = 3_500_000
output = 16_384
[modalities]
input = ["text", "image"]
output = ["text"]
[[weights]]
label = "Hugging Face"
url = "https://huggingface.co/meta-llama/Llama-4-Scout-17B-16E-Instruct"
+34
View File
@@ -0,0 +1,34 @@
name = "MiniMax-M2.1"
family = "minimax"
release_date = "2025-12-23"
last_updated = "2025-12-23"
attachment = false
reasoning = true
temperature = true
tool_call = true
open_weights = true
[limit]
context = 204_800
output = 131_072
[modalities]
input = ["text"]
output = ["text"]
[[weights]]
label = "Hugging Face"
url = "https://huggingface.co/MiniMaxAI/MiniMax-M2.1"
[[benchmarks]]
name = "SWE-Bench Verified"
score = 74
metric = "resolved"
source = "https://huggingface.co/MiniMaxAI/MiniMax-M2.1"
[[benchmarks]]
name = "SWE-Bench Pro"
score = 36.81
metric = "resolve rate"
dataset = "public"
source = "https://labs.scale.com/leaderboard/swe_bench_pro_public"
@@ -0,0 +1,21 @@
name = "MiniMax-M2.5-highspeed"
family = "minimax"
release_date = "2026-02-13"
last_updated = "2026-02-13"
attachment = false
reasoning = true
temperature = true
tool_call = true
open_weights = true
[limit]
context = 204_800
output = 131_072
[modalities]
input = ["text"]
output = ["text"]
[[weights]]
label = "Hugging Face"
url = "https://huggingface.co/MiniMaxAI/MiniMax-M2.5"
+48
View File
@@ -0,0 +1,48 @@
name = "MiniMax-M2.5"
family = "minimax"
release_date = "2026-02-12"
last_updated = "2026-02-12"
attachment = false
reasoning = true
temperature = true
tool_call = true
open_weights = true
[limit]
context = 204_800
output = 131_072
[modalities]
input = ["text"]
output = ["text"]
[[weights]]
label = "Hugging Face"
url = "https://huggingface.co/MiniMaxAI/MiniMax-M2.5"
[[benchmarks]]
name = "SWE-Bench Verified"
score = 75.8
metric = "resolved"
source = "https://www.swebench.com/"
[[benchmarks]]
name = "SWE-Atlas Codebase QnA"
score = 10.3
metric = "score"
harness = "Mini-SWE-Agent"
source = "https://labs.scale.com/leaderboard/sweatlas-qna"
[[benchmarks]]
name = "SWE-Atlas Refactoring"
score = 19.52
metric = "score"
harness = "Mini-SWE-Agent"
source = "https://labs.scale.com/leaderboard/sweatlas-refactoring"
[[benchmarks]]
name = "SWE-Atlas Test Writing"
score = 18.6
metric = "score"
harness = "Mini-SWE-Agent"
source = "https://labs.scale.com/leaderboard/sweatlas-tw"
@@ -0,0 +1,21 @@
name = "MiniMax-M2.7-highspeed"
family = "minimax"
release_date = "2026-03-18"
last_updated = "2026-03-18"
attachment = false
reasoning = true
temperature = true
tool_call = true
open_weights = true
[limit]
context = 204_800
output = 131_072
[modalities]
input = ["text"]
output = ["text"]
[[weights]]
label = "Hugging Face"
url = "https://huggingface.co/MiniMaxAI/MiniMax-M2.7"
+21
View File
@@ -0,0 +1,21 @@
name = "MiniMax-M2.7"
family = "minimax"
release_date = "2026-03-18"
last_updated = "2026-03-18"
attachment = false
reasoning = true
temperature = true
tool_call = true
open_weights = true
[limit]
context = 204_800
output = 131_072
[modalities]
input = ["text"]
output = ["text"]
[[weights]]
label = "Hugging Face"
url = "https://huggingface.co/MiniMaxAI/MiniMax-M2.7"
+27
View File
@@ -0,0 +1,27 @@
name = "MiniMax-M2"
family = "minimax"
release_date = "2025-10-27"
last_updated = "2025-10-27"
attachment = false
reasoning = true
temperature = true
tool_call = true
open_weights = true
[limit]
context = 196_608
output = 128_000
[modalities]
input = ["text"]
output = ["text"]
[[weights]]
label = "Hugging Face"
url = "https://huggingface.co/MiniMaxAI/MiniMax-M2"
[[benchmarks]]
name = "SWE-Bench Verified"
score = 69.4
metric = "resolved"
source = "https://huggingface.co/MiniMaxAI/MiniMax-M2"
+17
View File
@@ -0,0 +1,17 @@
name = "MiniMax-M3"
family = "minimax"
release_date = "2026-06-01"
last_updated = "2026-06-01"
attachment = true
reasoning = true
temperature = true
tool_call = true
open_weights = false
[limit]
context = 512_000
output = 128_000
[modalities]
input = ["text", "image", "video"]
output = ["text"]
+29
View File
@@ -0,0 +1,29 @@
name = "Codestral (latest)"
family = "codestral"
release_date = "2024-05-29"
last_updated = "2025-01-04"
attachment = false
reasoning = false
temperature = true
tool_call = true
knowledge = "2024-10"
open_weights = true
[limit]
context = 256_000
output = 4_096
[modalities]
input = ["text"]
output = ["text"]
[[weights]]
label = "Hugging Face"
url = "https://huggingface.co/mistralai/Codestral-22B-v0.1"
[[benchmarks]]
name = "Aider Polyglot"
score = 11.1
metric = "percent correct"
source = "https://aider.chat/docs/leaderboards/"
date = "2025-01-13"
+43
View File
@@ -0,0 +1,43 @@
name = "Devstral 2"
family = "devstral"
release_date = "2025-12-09"
last_updated = "2025-12-09"
attachment = false
reasoning = false
temperature = true
tool_call = true
knowledge = "2025-12"
open_weights = true
[limit]
context = 262_144
output = 262_144
[modalities]
input = ["text"]
output = ["text"]
[[weights]]
label = "Hugging Face"
url = "https://huggingface.co/mistralai/Devstral-2-123B-Instruct-2512"
[[benchmarks]]
name = "Artificial Analysis Coding Index"
score = 23.7
metric = "index"
source = "https://openrouter.ai/mistralai/devstral-2512/benchmarks"
date = "2026-05-31"
[[benchmarks]]
name = "SciCode"
score = 33.1
metric = "percent correct"
source = "https://openrouter.ai/mistralai/devstral-2512/benchmarks"
date = "2026-05-31"
[[benchmarks]]
name = "Terminal-Bench Hard"
score = 18.9
metric = "success rate"
source = "https://openrouter.ai/mistralai/devstral-2512/benchmarks"
date = "2026-05-31"
+25
View File
@@ -0,0 +1,25 @@
name = "Devstral Medium"
family = "devstral"
release_date = "2025-07-10"
last_updated = "2025-07-10"
attachment = false
reasoning = false
temperature = true
tool_call = true
knowledge = "2025-05"
open_weights = false
[limit]
context = 128_000
output = 128_000
[modalities]
input = ["text"]
output = ["text"]
[[benchmarks]]
name = "SWE-Bench Verified"
score = 61.6
metric = "resolved"
source = "https://mistral.ai/news/devstral-2507"
date = "2025-07-10"
@@ -0,0 +1,22 @@
name = "Devstral 2 (latest)"
family = "devstral"
release_date = "2025-12-02"
last_updated = "2025-12-02"
attachment = false
reasoning = false
temperature = true
tool_call = true
knowledge = "2025-12"
open_weights = true
[limit]
context = 262_144
output = 262_144
[modalities]
input = ["text"]
output = ["text"]
[[weights]]
label = "Hugging Face"
url = "https://huggingface.co/mistralai/Devstral-2-123B-Instruct-2512"
+29
View File
@@ -0,0 +1,29 @@
name = "Devstral Small"
family = "devstral"
release_date = "2025-07-10"
last_updated = "2025-07-10"
attachment = false
reasoning = false
temperature = true
tool_call = true
knowledge = "2025-05"
open_weights = true
[limit]
context = 128_000
output = 128_000
[modalities]
input = ["text"]
output = ["text"]
[[weights]]
label = "Hugging Face"
url = "https://huggingface.co/mistralai/Devstral-Small-2507"
[[benchmarks]]
name = "SWE-Bench Verified"
score = 53.6
metric = "resolved"
source = "https://mistral.ai/news/devstral-2507"
date = "2025-07-10"
@@ -0,0 +1,18 @@
name = "Magistral Medium (latest)"
family = "magistral-medium"
release_date = "2025-03-17"
last_updated = "2025-03-20"
attachment = false
reasoning = true
temperature = true
tool_call = true
knowledge = "2025-06"
open_weights = false
[limit]
context = 128_000
output = 16_384
[modalities]
input = ["text"]
output = ["text"]
+43
View File
@@ -0,0 +1,43 @@
name = "Mistral Large 2.1"
family = "mistral-large"
release_date = "2024-11-01"
last_updated = "2024-11-04"
attachment = false
reasoning = false
temperature = true
tool_call = true
knowledge = "2024-11"
open_weights = true
[limit]
context = 131_072
output = 16_384
[modalities]
input = ["text"]
output = ["text"]
[[weights]]
label = "Hugging Face"
url = "https://huggingface.co/mistralai/Mistral-Large-Instruct-2411"
[[benchmarks]]
name = "Artificial Analysis Coding Index"
score = 13.8
metric = "index"
source = "https://openrouter.ai/mistralai/mistral-large-2407/benchmarks"
date = "2026-03-11"
[[benchmarks]]
name = "SciCode"
score = 29.2
metric = "percent correct"
source = "https://openrouter.ai/mistralai/mistral-large-2407/benchmarks"
date = "2026-03-11"
[[benchmarks]]
name = "Terminal-Bench Hard"
score = 6.1
metric = "success rate"
source = "https://openrouter.ai/mistralai/mistral-large-2407/benchmarks"
date = "2026-03-11"

Some files were not shown because too many files have changed in this diff Show More