mirror of
https://github.com/BerriAI/litellm.git
synced 2026-09-14 23:21:35 +00:00
fix(cli): pin pi compat flags for gateway models
Co-Authored-By: Devin AI <158243242+devin-ai-integration[bot]@users.noreply.github.com>
This commit is contained in:
parent
67cb34ceee
commit
baa6446348
2 changed files with 4 additions and 0 deletions
|
|
@ -188,10 +188,13 @@ def provider_block(
|
|||
|
||||
Real contextWindow/maxTokens matter: pi otherwise assumes 128k/16384, which
|
||||
breaks compaction thresholds and over-asks models with smaller output caps.
|
||||
pi sniffs compat from the base URL, and one gateway URL fronts models with
|
||||
different capabilities, so both flags are pinned off.
|
||||
"""
|
||||
return { # mutable-ok: JSON serialization requires a mutable object
|
||||
"baseUrl": base_url.rstrip("/") + "/v1",
|
||||
"api": "openai-completions",
|
||||
"compat": {"supportsStore": False, "supportsLongCacheRetention": False},
|
||||
"apiKey": f"${LITELLM_PROXY_API_KEY_ENV}",
|
||||
"models": [_model_entry(model_id, limits) for model_id in model_ids], # mutable-ok: JSON array
|
||||
}
|
||||
|
|
|
|||
|
|
@ -188,6 +188,7 @@ class TestProviderBlock:
|
|||
assert block == {
|
||||
"baseUrl": "http://localhost:4000/v1",
|
||||
"api": "openai-completions",
|
||||
"compat": {"supportsStore": False, "supportsLongCacheRetention": False},
|
||||
"apiKey": "$LITELLM_PROXY_API_KEY",
|
||||
"models": [{"id": "m-1"}, {"id": "m-2"}],
|
||||
}
|
||||
|
|
|
|||
Loading…
Add table
Reference in a new issue