summaryrefslogtreecommitdiffstats
diff options
context:
space:
mode:
-rw-r--r--models.yaml587
1 files changed, 114 insertions, 473 deletions
diff --git a/models.yaml b/models.yaml
index 3f05eaa..2c70877 100644
--- a/models.yaml
+++ b/models.yaml
@@ -3,18 +3,11 @@
# - https://platform.openai.com/docs/api-reference/chat
- provider: openai
models:
- - name: gpt-5.1
+ - name: gpt-5.2
max_input_tokens: 400000
max_output_tokens: 128000
- input_price: 1.25
- output_price: 10
- supports_vision: true
- supports_function_calling: true
- - name: gpt-5.1-chat-latest
- max_input_tokens: 400000
- max_output_tokens: 128000
- input_price: 1.25
- output_price: 10
+ input_price: 1.75
+ output_price: 14
supports_vision: true
supports_function_calling: true
- name: gpt-5
@@ -24,13 +17,6 @@
output_price: 10
supports_vision: true
supports_function_calling: true
- - name: gpt-5-chat-latest
- max_input_tokens: 400000
- max_output_tokens: 128000
- input_price: 1.25
- output_price: 10
- supports_vision: true
- supports_function_calling: true
- name: gpt-5-mini
max_input_tokens: 400000
max_output_tokens: 128000
@@ -52,20 +38,6 @@
output_price: 8
supports_vision: true
supports_function_calling: true
- - name: gpt-4.1-mini
- max_input_tokens: 1047576
- max_output_tokens: 32768
- input_price: 0.4
- output_price: 1.6
- supports_vision: true
- supports_function_calling: true
- - name: gpt-4.1-nano
- max_input_tokens: 1047576
- max_output_tokens: 32768
- input_price: 0.1
- output_price: 0.4
- supports_vision: true
- supports_function_calling: true
- name: gpt-4o
max_input_tokens: 128000
max_output_tokens: 16384
@@ -73,91 +45,6 @@
output_price: 10
supports_vision: true
supports_function_calling: true
- - name: gpt-4o-mini
- max_input_tokens: 128000
- max_output_tokens: 16384
- input_price: 0.15
- output_price: 0.6
- supports_vision: true
- supports_function_calling: true
- - name: o4-mini
- max_input_tokens: 200000
- input_price: 1.1
- output_price: 4.4
- supports_vision: true
- supports_function_calling: true
- system_prompt_prefix: Formatting re-enabled
- patch:
- body:
- max_tokens: null
- temperature: null
- top_p: null
- - name: o4-mini-high
- real_name: o4-mini
- max_input_tokens: 200000
- input_price: 1.1
- output_price: 4.4
- supports_vision: true
- supports_function_calling: true
- system_prompt_prefix: Formatting re-enabled
- patch:
- body:
- reasoning_effort: high
- max_tokens: null
- temperature: null
- top_p: null
- - name: o3
- max_input_tokens: 200000
- input_price: 2
- output_price: 8
- supports_vision: true
- supports_function_calling: true
- system_prompt_prefix: Formatting re-enabled
- patch:
- body:
- max_tokens: null
- temperature: null
- top_p: null
- - name: o3-high
- real_name: o3
- max_input_tokens: 200000
- input_price: 2
- output_price: 8
- supports_vision: true
- supports_function_calling: true
- system_prompt_prefix: Formatting re-enabled
- patch:
- body:
- reasoning_effort: high
- max_tokens: null
- temperature: null
- top_p: null
- - name: o3-mini
- max_input_tokens: 200000
- input_price: 1.1
- output_price: 4.4
- supports_vision: true
- supports_function_calling: true
- system_prompt_prefix: Formatting re-enabled
- patch:
- body:
- max_tokens: null
- temperature: null
- top_p: null
- - name: o3-mini-high
- real_name: o3-mini
- max_input_tokens: 200000
- input_price: 1.1
- output_price: 4.4
- supports_vision: true
- supports_function_calling: true
- system_prompt_prefix: Formatting re-enabled
- patch:
- body:
- reasoning_effort: high
- max_tokens: null
- temperature: null
- top_p: null
- name: gpt-4-turbo
max_input_tokens: 128000
max_output_tokens: 4096
@@ -211,6 +98,14 @@
output_price: 0
supports_vision: true
supports_function_calling: true
+ - name: gemini-3-pro-preview
+ max_input_tokens: 1048576
+ supports_vision: true
+ supports_function_calling: true
+ - name: gemini-3-flash-preview
+ max_input_tokens: 1048576
+ supports_vision: true
+ supports_function_calling: true
- name: gemini-2.0-flash
max_input_tokens: 1048576
max_output_tokens: 8192
@@ -400,6 +295,12 @@
# - https://docs.mistral.ai/api/
- provider: mistral
models:
+ - name: mistral-large-latest
+ max_output_tokens: 262144
+ input_price: 0.5
+ output_price: 1.5
+ supports_function_calling: true
+ supports_vision: true
- name: mistral-medium-latest
max_input_tokens: 131072
input_price: 0.4
@@ -413,28 +314,33 @@
supports_function_calling: true
supports_vision: true
- name: magistral-medium-latest
- max_input_tokens: 40960
+ max_input_tokens: 131072
input_price: 2
output_price: 5
- name: magistral-small-latest
- max_input_tokens: 40960
+ max_input_tokens: 131072
input_price: 0.5
output_price: 1.5
- name: devstral-medium-latest
- max_input_tokens: 256000
+ max_input_tokens: 262144
input_price: 0.4
output_price: 2
supports_function_calling: true
- name: devstral-small-latest
- max_input_tokens: 256000
+ max_input_tokens: 262144
input_price: 0.1
output_price: 0.3
supports_function_calling: true
- name: codestral-latest
- max_input_tokens: 256000
+ max_input_tokens: 262144
input_price: 0.3
output_price: 0.9
supports_function_calling: true
+ - name: ministral-14b-latest
+ max_input_tokens: 262144
+ input_price: 0.2
+ output_price: 0.2
+ supports_function_calling: true
- name: mistral-embed
type: embedding
max_input_tokens: 8092
@@ -520,36 +426,21 @@
# - https://docs.x.ai/docs/api-reference#chat-completions
- provider: xai
models:
- - name: grok-4
- max_input_tokens: 256000
- input_price: 3
- output_price: 15
- supports_function_calling: true
- - name: grok-4-fast-non-reasoning
+ - name: grok-4-1-fast-non-reasoning
max_input_tokens: 2000000
input_price: 0.2
output_price: 0.5
supports_function_calling: true
- - name: grok-4-fast-reasoning
+ - name: grok-4-1-fast-reasoning
max_input_tokens: 2000000
input_price: 0.2
output_price: 0.5
supports_function_calling: true
- - name: grok-code-fast
+ - name: grok-code-fast-1
max_input_tokens: 256000
input_price: 0.2
output_price: 1.5
supports_function_calling: true
- - name: grok-3
- max_input_tokens: 131072
- input_price: 3
- output_price: 15
- supports_function_calling: true
- - name: grok-3-mini
- max_input_tokens: 131072
- input_price: 0.3
- output_price: 0.5
- supports_function_calling: true
# Links:
# - https://docs.perplexity.ai/getting-started/models
@@ -568,10 +459,6 @@
max_input_tokens: 128000
input_price: 2
output_price: 8
- - name: sonar-reasoning
- max_input_tokens: 128000
- input_price: 1
- output_price: 5
- name: sonar-deep-research
max_input_tokens: 128000
input_price: 2
@@ -654,6 +541,14 @@
output_price: 0.4
supports_vision: true
supports_function_calling: true
+ - name: gemini-3-pro-preview
+ max_input_tokens: 1048576
+ supports_vision: true
+ supports_function_calling: true
+ - name: gemini-3-flash-preview
+ max_input_tokens: 1048576
+ supports_vision: true
+ supports_function_calling: true
- name: gemini-2.0-flash-001
max_input_tokens: 1048576
max_output_tokens: 8192
@@ -814,16 +709,6 @@
output_price: 4
supports_vision: true
supports_function_calling: true
- - name: mistral-small-2503
- max_input_tokens: 32000
- input_price: 0.1
- output_price: 0.3
- supports_function_calling: true
- - name: codestral-2501
- max_input_tokens: 256000
- input_price: 0.3
- output_price: 0.9
- supports_function_calling: true
- name: text-embedding-005
type: embedding
max_input_tokens: 20000
@@ -1246,27 +1131,20 @@
# - https://cloud.tencent.com/document/product/1729/111007
- provider: hunyuan
models:
- - name: hunyuan-turbos-latest
- max_input_tokens: 28000
+ - name: hunyuan-2.0-instruct-20251111
+ max_input_tokens: 131072
input_price: 0.112
output_price: 0.28
supports_function_calling: true
- - name: hunyuan-t1-latest
- max_input_tokens: 28000
+ - name: hunyuan-2.0-thinking-20251109
+ max_input_tokens: 131072
input_price: 0.14
output_price: 0.56
- - name: hunyuan-lite
- max_input_tokens: 250000
- input_price: 0
- output_price: 0
supports_function_calling: true
- - name: hunyuan-turbos-vision
- max_input_tokens: 6144
+ - name: hunyuan-vision-1.5-instruct
+ max_input_tokens: 24576
input_price: 0.42
- output_price: 0.84
- supports_vision: true
- - name: hunyuan-t1-vision
- max_input_tokens: 24000
+ output_price: 1.26
supports_vision: true
- name: hunyuan-embedding
type: embedding
@@ -1325,54 +1203,27 @@
# - https://open.bigmodel.cn/dev/api#glm-4
- provider: zhipuai
models:
- - name: glm-4.6
+ - name: glm-4.7
max_input_tokens: 202752
- input_price: 0.28
- output_price: 1.12
- supports_function_calling: true
- - name: glm-4.5
- max_input_tokens: 131072
- input_price: 0.28
- output_price: 1.12
- - name: glm-4.5-x
- max_input_tokens: 131072
- input_price: 1.12
- output_price: 4.48
+ input_price: 0.56
+ output_price: 2.24
supports_function_calling: true
- - name: glm-4.5-air
- max_input_tokens: 131072
- input_price: 0.084
- output_price: 0.56
- - name: glm-4.5-airx
- max_input_tokens: 131072
+
+ - name: glm-4.7:instruct
+ real_name: glm-4.7
+ max_input_tokens: 202752
input_price: 0.56
output_price: 2.24
supports_function_calling: true
- - name: glm-4.5-flash
- max_input_tokens: 131072
- input_price: 0
- output_price: 0
- - name: glm-4.5v
+ patch:
+ body:
+ thinking:
+ type: disabled
+ - name: glm-4.6v
max_input_tokens: 65536
- input_price: 0.56
- output_price: 1.68
+ input_price: 0.28
+ output_price: 0.84
supports_vision: true
- - name: glm-z1-air
- max_input_tokens: 131072
- input_price: 0.07
- output_price: 0.07
- - name: glm-z1-airx
- max_input_tokens: 131072
- input_price: 0.7
- output_price: 0.7
- - name: glm-z1-flashx
- max_input_tokens: 131072
- input_price: 0.014
- output_price: 0.014
- - name: glm-z1-flash
- max_input_tokens: 131072
- input_price: 0
- output_price: 0
- name: embedding-3
type: embedding
max_input_tokens: 8192
@@ -1389,29 +1240,27 @@
# - https://platform.minimaxi.com/document/ChatCompletion%20v2
- provider: minimax
models:
- - name: minimax-m2
+ - name: minimax-m2.1
max_input_tokens: 204800
input_price: 0.294
output_price: 1.176
supports_function_calling: true
+ - name: minimax-m2.1-lightning
+ max_input_tokens: 204800
+ input_price: 0.294
+ output_price: 2.352
+ supports_function_calling: true
# Links:
# - https://openrouter.ai/models
# - https://openrouter.ai/docs/api-reference/chat-completion
- provider: openrouter
models:
- - name: openai/gpt-5.1
+ - name: openai/gpt-5.2
max_input_tokens: 400000
max_output_tokens: 128000
- input_price: 1.25
- output_price: 10
- supports_vision: true
- supports_function_calling: true
- - name: openai/gpt-5.1-chat
- max_input_tokens: 400000
- max_output_tokens: 128000
- input_price: 1.25
- output_price: 10
+ input_price: 1.75
+ output_price: 14
supports_vision: true
supports_function_calling: true
- name: openai/gpt-5
@@ -1421,13 +1270,6 @@
output_price: 10
supports_vision: true
supports_function_calling: true
- - name: openai/gpt-5-chat
- max_input_tokens: 400000
- max_output_tokens: 128000
- input_price: 1.25
- output_price: 10
- supports_vision: true
- supports_function_calling: true
- name: openai/gpt-5-mini
max_input_tokens: 400000
max_output_tokens: 128000
@@ -1449,104 +1291,12 @@
output_price: 8
supports_vision: true
supports_function_calling: true
- - name: openai/gpt-4.1-mini
- max_input_tokens: 1047576
- max_output_tokens: 32768
- input_price: 0.4
- output_price: 1.6
- supports_vision: true
- supports_function_calling: true
- - name: openai/gpt-4.1-nano
- max_input_tokens: 1047576
- max_output_tokens: 32768
- input_price: 0.1
- output_price: 0.4
- supports_vision: true
- supports_function_calling: true
- name: openai/gpt-4o
max_input_tokens: 128000
input_price: 2.5
output_price: 10
supports_vision: true
supports_function_calling: true
- - name: openai/gpt-4o-mini
- max_input_tokens: 128000
- input_price: 0.15
- output_price: 0.6
- supports_vision: true
- supports_function_calling: true
- - name: openai/o4-mini
- max_input_tokens: 200000
- input_price: 1.1
- output_price: 4.4
- supports_vision: true
- supports_function_calling: true
- system_prompt_prefix: Formatting re-enabled
- patch:
- body:
- max_tokens: null
- temperature: null
- top_p: null
- - name: openai/o4-mini-high
- max_input_tokens: 200000
- input_price: 1.1
- output_price: 4.4
- supports_vision: true
- supports_function_calling: true
- system_prompt_prefix: Formatting re-enabled
- patch:
- body:
- reasoning_effort: high
- max_tokens: null
- temperature: null
- top_p: null
- - name: openai/o3
- max_input_tokens: 200000
- input_price: 2
- output_price: 8
- supports_vision: true
- supports_function_calling: true
- system_prompt_prefix: Formatting re-enabled
- patch:
- body:
- max_tokens: null
- temperature: null
- top_p: null
- - name: openai/o3-high
- real_name: openai/o3
- max_input_tokens: 200000
- input_price: 2
- output_price: 8
- supports_vision: true
- supports_function_calling: true
- system_prompt_prefix: Formatting re-enabled
- patch:
- body:
- reasoning_effort: high
- temperature: null
- top_p: null
- - name: openai/o3-mini
- max_input_tokens: 200000
- input_price: 1.1
- output_price: 4.4
- supports_vision: true
- supports_function_calling: true
- system_prompt_prefix: Formatting re-enabled
- patch:
- body:
- temperature: null
- top_p: null
- - name: openai/o3-mini-high
- max_input_tokens: 200000
- input_price: 1.1
- output_price: 4.4
- supports_vision: true
- supports_function_calling: true
- system_prompt_prefix: Formatting re-enabled
- patch:
- body:
- temperature: null
- top_p: null
- name: openai/gpt-oss-120b
max_input_tokens: 131072
input_price: 0.09
@@ -1662,6 +1412,11 @@
max_input_tokens: 131072
input_price: 0.12
output_price: 0.3
+ - name: mistralai/mistral-large-2512
+ max_input_tokens: 262144
+ input_price: 0.5
+ output_price: 1.5
+ supports_function_calling: true
- name: mistralai/mistral-medium-3.1
max_input_tokens: 131072
input_price: 0.4
@@ -1673,22 +1428,10 @@
input_price: 0.1
output_price: 0.3
supports_vision: true
- - name: mistralai/magistral-medium-2506
- max_input_tokens: 40960
- input_price: 2
- output_price: 5
- - name: mistralai/magistral-medium-2506:thinking
- max_input_tokens: 40960
- input_price: 2
- output_price: 5
- - name: mistralai/magistral-small-2506
- max_input_tokens: 40960
+ - name: mistralai/devstral-2512
+ max_input_tokens: 262144
input_price: 0.5
- output_price: 1.5
- - name: mistralai/devstral-medium
- max_input_tokens: 131072
- input_price: 0.4
- output_price: 2
+ output_price: 0.22
supports_function_calling: true
- name: mistralai/devstral-small
max_input_tokens: 131072
@@ -1700,6 +1443,11 @@
input_price: 0.3
output_price: 0.9
supports_function_calling: true
+ - name: mistralai/ministral-14b-2512
+ max_input_tokens: 262144
+ input_price: 0.2
+ output_price: 0.2
+ supports_function_calling: true
- name: ai21/jamba-large-1.7
max_input_tokens: 256000
input_price: 2
@@ -1720,25 +1468,10 @@
max_output_tokens: 4096
input_price: 0.0375
output_price: 0.15
- - name: deepseek/deepseek-v3.2-exp
- max_input_tokens: 163840
- input_price: 0.27
- output_price: 0.40
- - name: deepseek/deepseek-v3.1-terminus
- max_input_tokens: 163840
- input_price: 0.23
- output_price: 0.90
- - name: deepseek/deepseek-chat-v3.1
+ - name: deepseek/deepseek-v3.2
max_input_tokens: 163840
- input_price: 0.2
- output_price: 0.8
- - name: deepseek/deepseek-r1-0528
- max_input_tokens: 128000
- input_price: 0.50
- output_price: 2.15
- patch:
- body:
- include_reasoning: true
+ input_price: 0.25
+ output_price: 0.38
- name: qwen/qwen3-max
max_input_tokens: 262144
input_price: 1.2
@@ -1821,12 +1554,7 @@
input_price: 0.29
output_price: 1.15
supports_function_calling: true
- - name: x-ai/grok-4
- max_input_tokens: 256000
- input_price: 3
- output_price: 15
- supports_function_calling: true
- - name: x-ai/grok-4-fast
+ - name: x-ai/grok-4.1-fast
max_input_tokens: 2000000
input_price: 0.2
output_price: 0.5
@@ -1873,13 +1601,6 @@
patch:
body:
include_reasoning: true
- - name: perplexity/sonar-reasoning
- max_input_tokens: 127000
- input_price: 1
- output_price: 5
- patch:
- body:
- include_reasoning: true
- name: perplexity/sonar-deep-research
max_input_tokens: 200000
input_price: 2
@@ -1887,15 +1608,21 @@
patch:
body:
include_reasoning: true
- - name: minimax/minimax-m2
+ - name: minimax/minimax-m2.1
max_input_tokens: 196608
- input_price: 0.15
- output_price: 0.45
- - name: z-ai/glm-4.6
+ input_price: 0.12
+ output_price: 0.48
+ supports_function_calling: true
+ - name: z-ai/glm-4.7
max_input_tokens: 202752
- input_price: 0.5
- output_price: 1.75
+ input_price: 0.16
+ output_price: 0.80
supports_function_calling: true
+ - name: z-ai/glm-4.6v
+ max_input_tokens: 131072
+ input_price: 0.3
+ output_price: 0.9
+ supports_vision: true
# Links:
# - https://github.com/marketplace?type=models
@@ -1906,11 +1633,6 @@
max_output_tokens: 128000
supports_vision: true
supports_function_calling: true
- - name: gpt-5-chat
- max_input_tokens: 400000
- max_output_tokens: 128000
- supports_vision: true
- supports_function_calling: true
- name: gpt-5-mini
max_input_tokens: 400000
max_output_tokens: 128000
@@ -1926,90 +1648,10 @@
max_output_tokens: 32768
supports_vision: true
supports_function_calling: true
- - name: gpt-4.1-mini
- max_input_tokens: 1047576
- max_output_tokens: 32768
- supports_vision: true
- supports_function_calling: true
- - name: gpt-4.1-nano
- max_input_tokens: 1047576
- max_output_tokens: 32768
- supports_vision: true
- supports_function_calling: true
- name: gpt-4o
max_input_tokens: 128000
max_output_tokens: 16384
supports_function_calling: true
- - name: gpt-4o-mini
- max_input_tokens: 128000
- max_output_tokens: 16384
- supports_function_calling: true
- - name: o4-mini
- max_input_tokens: 200000
- supports_vision: true
- supports_function_calling: true
- system_prompt_prefix: Formatting re-enabled
- patch:
- body:
- max_tokens: null
- temperature: null
- top_p: null
- - name: o4-mini-high
- real_name: o4-mini
- max_input_tokens: 200000
- supports_vision: true
- supports_function_calling: true
- system_prompt_prefix: Formatting re-enabled
- patch:
- body:
- reasoning_effort: high
- max_tokens: null
- temperature: null
- top_p: null
- - name: o3
- max_input_tokens: 200000
- supports_vision: true
- supports_function_calling: true
- system_prompt_prefix: Formatting re-enabled
- patch:
- body:
- max_tokens: null
- temperature: null
- top_p: null
- - name: o3-high
- real_name: o3
- max_input_tokens: 200000
- supports_vision: true
- supports_function_calling: true
- system_prompt_prefix: Formatting re-enabled
- patch:
- body:
- reasoning_effort: high
- max_tokens: null
- temperature: null
- top_p: null
- - name: o3-mini
- max_input_tokens: 200000
- supports_vision: true
- supports_function_calling: true
- system_prompt_prefix: Formatting re-enabled
- patch:
- body:
- max_tokens: null
- temperature: null
- top_p: null
- - name: o3-mini-high
- real_name: o3-mini
- max_input_tokens: 200000
- supports_vision: true
- supports_function_calling: true
- system_prompt_prefix: Formatting re-enabled
- patch:
- body:
- reasoning_effort: high
- max_tokens: null
- temperature: null
- top_p: null
- name: text-embedding-3-large
type: embedding
max_tokens_per_chunk: 8191
@@ -2128,22 +1770,11 @@
input_price: 0.18
output_price: 0.69
supports_vision: true
- - name: deepseek-ai/DeepSeek-V3.2-Exp
- max_input_tokens: 163840
- input_price: 0.27
- output_price: 0.40
- - name: deepseek-ai/DeepSeek-V3.1-Terminus
- max_input_tokens: 163840
- input_price: 0.27
- output_price: 1.0
- - name: deepseek-ai/DeepSeek-V3.1
- max_input_tokens: 163840
- input_price: 0.3
- output_price: 1.0
- - name: deepseek-ai/DeepSeek-R1-0528
+ - name: deepseek-ai/DeepSeek-V3.2
max_input_tokens: 163840
- input_price: 0.5
- output_price: 2.15
+ input_price: 0.26
+ output_price: 0.39
+ supports_function_calling: true
- name: google/gemma-3-27b-it
max_input_tokens: 131072
input_price: 0.1
@@ -2162,11 +1793,21 @@
input_price: 0.55
output_price: 2.5
supports_function_calling: true
- - name: zai-org/GLM-4.6
+ - name: MiniMaxAI/MiniMax-M2.1
+ max_input_tokens: 262144
+ input_price: 0.28
+ output_price: 1.2
+ supports_function_calling: true
+ - name: zai-org/GLM-4.7
max_input_tokens: 202752
- input_price: 0.6
- output_price: 1.9
+ input_price: 0.43
+ output_price: 1.75
supports_function_calling: true
+ - name: zai-org/GLM-4.6V
+ max_input_tokens: 131072
+ input_price: 0.3
+ output_price: 0.9
+ supports_vision: true
- name: BAAI/bge-large-en-v1.5
type: embedding
input_price: 0.01