mirror of
https://github.com/BerriAI/litellm.git
synced 2026-10-02 02:11:58 +00:00
chore(dashscope): add qwen3-rerank context-window limits
Fill in the context-window metadata other rerank entries carry, for dashscope/qwen3-rerank in both price JSONs: request max input (120k), single query/document max (4k), and max document count (500). Values from the Aliyun Model Studio text-rerank docs. Co-Authored-By: Claude Opus 4.7 (1M context) <noreply@anthropic.com>
This commit is contained in:
parent
8acc9b9f3b
commit
4bb567f976
2 changed files with 12 additions and 0 deletions
|
|
@ -10995,6 +10995,12 @@
|
|||
"dashscope/qwen3-rerank": {
|
||||
"input_cost_per_token": 1e-07,
|
||||
"litellm_provider": "dashscope",
|
||||
"max_document_chunks_per_query": 500,
|
||||
"max_input_tokens": 120000,
|
||||
"max_output_tokens": 120000,
|
||||
"max_query_tokens": 4000,
|
||||
"max_tokens": 120000,
|
||||
"max_tokens_per_document_chunk": 4000,
|
||||
"mode": "rerank",
|
||||
"output_cost_per_token": 0.0,
|
||||
"source": "https://www.alibabacloud.com/help/en/model-studio/model-pricing"
|
||||
|
|
|
|||
|
|
@ -10995,6 +10995,12 @@
|
|||
"dashscope/qwen3-rerank": {
|
||||
"input_cost_per_token": 1e-07,
|
||||
"litellm_provider": "dashscope",
|
||||
"max_document_chunks_per_query": 500,
|
||||
"max_input_tokens": 120000,
|
||||
"max_output_tokens": 120000,
|
||||
"max_query_tokens": 4000,
|
||||
"max_tokens": 120000,
|
||||
"max_tokens_per_document_chunk": 4000,
|
||||
"mode": "rerank",
|
||||
"output_cost_per_token": 0.0,
|
||||
"source": "https://www.alibabacloud.com/help/en/model-studio/model-pricing"
|
||||
|
|
|
|||
Loading…
Add table
Reference in a new issue