diff options
| author | sigoden <sigoden@gmail.com> | 2025-02-17 08:45:21 +0800 |
|---|---|---|
| committer | GitHub <noreply@github.com> | 2025-02-17 08:45:21 +0800 |
| commit | a6eda31de56044319e0207a4ffc1869e26224b26 (patch) | |
| tree | 1d288dbfad92546305d03a0a1294d2cced450658 /models.yaml | |
| parent | f14f1ad01bf31bb4bb4322661691be9392c30acf (diff) | |
| download | aichat-a6eda31de56044319e0207a4ffc1869e26224b26.tar.gz | |
feat: remove supports for fireworks/siliconflow/together (#1181)
Diffstat (limited to 'models.yaml')
| -rw-r--r-- | models.yaml | 288 |
1 files changed, 0 insertions, 288 deletions
diff --git a/models.yaml b/models.yaml index e3bd1b2..217a716 100644 --- a/models.yaml +++ b/models.yaml @@ -1695,294 +1695,6 @@ max_batch_size: 100 # Links: -# - https://fireworks.ai/models -# - https://fireworks.ai/pricing -# - https://docs.fireworks.ai/api-reference/post-chatcompletions -- provider: fireworks - models: - - name: accounts/fireworks/models/llama-v3p3-70b-instruct - max_input_tokens: 131072 - input_price: 0.9 - output_price: 0.9 - - name: accounts/fireworks/models/llama-v3p1-405b-instruct - max_input_tokens: 131072 - input_price: 3 - output_price: 3 - supports_function_calling: true - - name: accounts/fireworks/models/llama-v3p1-70b-instruct - max_input_tokens: 131072 - input_price: 0.9 - output_price: 0.9 - supports_function_calling: true - - name: accounts/fireworks/models/llama-v3p1-8b-instruct - max_input_tokens: 131072 - input_price: 0.2 - output_price: 0.2 - - name: accounts/fireworks/models/llama-v3p2-90b-vision-instruct - max_input_tokens: 131072 - input_price: 0.9 - output_price: 0.9 - supports_vision: true - - name: accounts/fireworks/models/llama-v3p2-11b-vision-instruct - max_input_tokens: 131072 - input_price: 0.2 - output_price: 0.2 - supports_vision: true - - name: accounts/fireworks/models/qwen2p5-72b-instruct - max_input_tokens: 32768 - input_price: 0.9 - output_price: 0.9 - supports_function_calling: true - - name: accounts/fireworks/models/qwen2p5-coder-32b-instruct - max_input_tokens: 32768 - input_price: 0.9 - output_price: 0.9 - - name: accounts/fireworks/models/qwen-qwq-32b-preview - max_input_tokens: 32768 - input_price: 0.9 - output_price: 0.9 - - name: accounts/fireworks/models/qwen2-vl-72b-instruct - max_input_tokens: 32768 - input_price: 0.9 - output_price: 0.9 - supports_vision: true - - name: accounts/fireworks/models/deepseek-v3 - max_input_tokens: 131072 - input_price: 0.9 - output_price: 0.9 - - name: accounts/fireworks/models/deepseek-r1 - max_input_tokens: 160000 - input_price: 8 - output_price: 8 - - name: accounts/fireworks/models/deepseek-r1-distill-llama-70b - max_input_tokens: 131072 - input_price: 0.9 - output_price: 0.9 - - name: accounts/fireworks/models/deepseek-r1-distill-qwen-32b - max_input_tokens: 131072 - input_price: 0.9 - output_price: 0.9 - - name: accounts/fireworks/models/mistral-small-24b-instruct-2501 - max_input_tokens: 32768 - input_price: 0.9 - output_price: 0.9 - - name: nomic-ai/nomic-embed-text-v1.5 - type: embedding - input_price: 0.008 - max_tokens_per_chunk: 8192 - default_chunk_size: 1500 - max_batch_size: 100 - - name: WhereIsAI/UAE-Large-V1 - type: embedding - input_price: 0.016 - max_tokens_per_chunk: 512 - default_chunk_size: 1000 - max_batch_size: 100 - - name: thenlper/gte-large - type: embedding - input_price: 0.016 - max_tokens_per_chunk: 512 - default_chunk_size: 1000 - max_batch_size: 100 - -# Links -# - https://cloud.siliconflow.cn/models -# - https://docs.siliconflow.cn/api-reference/chat-completions/chat-completions -- provider: siliconflow - models: - - name: meta-llama/Llama-3.3-70B-Instruct - max_input_tokens: 32768 - input_price: 0.578 - output_price: 0.578 - - name: meta-llama/Meta-Llama-3.1-405B-Instruct - max_input_tokens: 32768 - input_price: 2.94 - output_price: 2.94 - - name: meta-llama/Meta-Llama-3.1-70B-Instruct - max_input_tokens: 32768 - input_price: 0.578 - output_price: 0.578 - - name: meta-llama/Meta-Llama-3.1-8B-Instruct - max_input_tokens: 32768 - input_price: 0 - output_price: 0 - - name: Qwen/Qwen2.5-72B-Instruct - max_input_tokens: 32768 - input_price: 0.578 - output_price: 0.578 - supports_function_calling: true - - name: Qwen/Qwen2.5-72B-Instruct-128K - max_input_tokens: 131072 - input_price: 0.578 - output_price: 0.578 - supports_function_calling: true - - name: Qwen/Qwen2.5-7B-Instruct - max_input_tokens: 32768 - input_price: 0 - output_price: 0 - supports_function_calling: true - - name: Qwen/Qwen2.5-Coder-32B-Instruct - max_input_tokens: 32768 - input_price: 0.176 - output_price: 0.176 - - name: Qwen/Qwen2.5-Coder-7B-Instruct - max_input_tokens: 32768 - input_price: 0 - output_price: 0 - - name: Qwen/Qwen2-VL-72B-Instruct - max_input_tokens: 32768 - input_price: 0.5782 - output_price: 0.5782 - supports_vision: true - - name: Qwen/QVQ-72B-Preview - max_input_tokens: 32768 - input_price: 1.386 - output_price: 1.386 - supports_vision: true - - name: Qwen/QwQ-32B-Preview - max_input_tokens: 32768 - input_price: 0.176 - output_price: 0.176 - - name: deepseek-ai/DeepSeek-V3 - max_input_tokens: 65536 - input_price: 0.28 - output_price: 0.28 - - name: deepseek-ai/DeepSeek-R1 - max_input_tokens: 65536 - input_price: 2.24 - output_price: 2.24 - - name: deepseek-ai/DeepSeek-R1-Distill-Llama-70B - max_input_tokens: 32768 - input_price: 0.578 - output_price: 0.578 - - name: deepseek-ai/DeepSeek-R1-Distill-Qwen-32B - max_input_tokens: 32768 - input_price: 0.176 - output_price: 0.176 - - name: deepseek-ai/DeepSeek-V2.5 - max_input_tokens: 32768 - input_price: 0.7 - output_price: 0.7 - supports_function_calling: true - - name: deepseek-ai/deepseek-vl2 - max_input_tokens: 32768 - input_price: 0.138 - output_price: 0.138 - supports_vision: true - - name: BAAI/bge-large-en-v1.5 - type: embedding - input_price: 0 - max_tokens_per_chunk: 512 - default_chunk_size: 1000 - max_batch_size: 100 - - name: BAAI/bge-large-zh-v1.5 - type: embedding - input_price: 0 - max_tokens_per_chunk: 512 - default_chunk_size: 1000 - max_batch_size: 100 - - name: BAAI/bge-m3 - type: embedding - input_price: 0 - max_tokens_per_chunk: 8192 - default_chunk_size: 2000 - max_batch_size: 100 - - name: BAAI/bge-reranker-v2-m3 - type: reranker - max_input_tokens: 8192 - input_price: 0 - -# Links: -# - https://docs.together.ai/docs/serverless-models -# - https://www.together.ai/pricing -# - https://docs.together.ai/reference/chat-completions-1 -- provider: together - models: - - name: meta-llama/Llama-3.3-70B-Instruct-Turbo - max_input_tokens: 131072 - input_price: 0.88 - output_price: 0.88 - supports_function_calling: true - - name: meta-llama/Meta-Llama-3.1-405B-Instruct-Turbo - max_input_tokens: 130815 - input_price: 3.5 - output_price: 3.5 - supports_function_calling: true - - name: meta-llama/Meta-Llama-3.1-70B-Instruct-Turbo - max_input_tokens: 131072 - input_price: 0.88 - output_price: 0.88 - supports_function_calling: true - - name: meta-llama/Meta-Llama-3.1-8B-Instruct-Turbo - max_input_tokens: 131072 - input_price: 0.18 - output_price: 0.18 - supports_function_calling: true - - name: meta-llama/Llama-3.2-90B-Vision-Instruct-Turbo - max_input_tokens: 131072 - input_price: 0.88 - output_price: 0.88 - supports_vision: true - - name: meta-llama/Llama-3.2-11B-Vision-Instruct-Turbo - max_input_tokens: 131072 - input_price: 0.18 - output_price: 0.18 - supports_vision: true - - name: Qwen/Qwen2.5-72B-Instruct-Turbo - max_input_tokens: 32768 - input_price: 1.2 - output_price: 1.2 - - name: Qwen/Qwen2.5-7B-Instruct-Turbo - max_input_tokens: 32768 - input_price: 0.3 - output_price: 0.3 - - name: Qwen/Qwen2.5-Coder-32B-Instruct - max_input_tokens: 32768 - input_price: 0.8 - output_price: 0.8 - - name: Qwen/QwQ-32B-Preview - max_input_tokens: 32768 - input_price: 1.2 - output_price: 1.2 - - name: Qwen/Qwen2-VL-72B-Instruct - max_input_tokens: 32768 - input_price: 1.2 - output_price: 1.2 - supports_vision: true - - name: deepseek-ai/DeepSeek-V3 - max_input_tokens: 131072 - input_price: 1.25 - output_price: 1.25 - - name: deepseek-ai/DeepSeek-R1 - max_input_tokens: 163840 - input_price: 7 - output_price: 7 - - name: deepseek-ai/DeepSeek-R1-Distill-Llama-70B - max_input_tokens: 131072 - input_price: 2 - output_price: 2 - - name: mistralai/Mistral-Small-24B-Instruct-2501 - max_input_tokens: 32768 - input_price: 0.8 - output_price: 0.8 - - name: WhereIsAI/UAE-Large-V1 - type: embedding - input_price: 0.016 - max_tokens_per_chunk: 512 - default_chunk_size: 1000 - max_batch_size: 100 - - name: BAAI/bge-large-en-v1.5 - type: embedding - input_price: 0.016 - max_tokens_per_chunk: 512 - default_chunk_size: 1000 - max_batch_size: 100 - - name: Salesforce/Llama-Rank-V1 - type: reranker - max_input_tokens: 8192 - input_price: 0.1 - -# Links: # - https://jina.ai/models # - https://api.jina.ai/redoc - provider: jina |
