fix(proxy): clamp parallel slot TTL override to a positive value

Co-Authored-By: Devin AI <158243242+devin-ai-integration[bot]@users.noreply.github.com>
This commit is contained in:
milan 2026-08-20 01:05:36 +00:00
parent 66dbf2101c
commit 4e608c5426
2 changed files with 2 additions and 2 deletions

View file

@ -369,7 +369,7 @@ PROJECT_OTPM_DESCRIPTOR_KEY: Final = "model_per_project_otpm"
# considered leaked (worker crashed without any release callback firing) and
# pruned. Also the longest request duration the gauge can track: a request
# running longer than this stops occupying its slot.
PARALLEL_REQUEST_SLOT_TTL_SECONDS: Final = int(os.getenv("LITELLM_PARALLEL_REQUEST_SLOT_TTL_SECONDS", "3600"))
PARALLEL_REQUEST_SLOT_TTL_SECONDS: Final = max(1, int(os.getenv("LITELLM_PARALLEL_REQUEST_SLOT_TTL_SECONDS", "3600")))
CacheCounterValue: TypeAlias = int | float | str | bytes

View file

@ -4960,7 +4960,7 @@ def test_pre_call_redacts_and_masks_raw_request(logging_obj):
assert "key=*****" in raw_api_base
def _streaming_logging_obj_with_callbacks(callbacks):
def _streaming_logging_obj_with_callbacks(callbacks: list[CustomLogger]):
from datetime import datetime
obj = LitellmLogging(