mirror of
https://github.com/BerriAI/litellm.git
synced 2026-10-07 02:59:05 +00:00
fix(proxy): clamp parallel slot TTL override to a positive value
Co-Authored-By: Devin AI <158243242+devin-ai-integration[bot]@users.noreply.github.com>
This commit is contained in:
parent
66dbf2101c
commit
4e608c5426
2 changed files with 2 additions and 2 deletions
|
|
@ -369,7 +369,7 @@ PROJECT_OTPM_DESCRIPTOR_KEY: Final = "model_per_project_otpm"
|
|||
# considered leaked (worker crashed without any release callback firing) and
|
||||
# pruned. Also the longest request duration the gauge can track: a request
|
||||
# running longer than this stops occupying its slot.
|
||||
PARALLEL_REQUEST_SLOT_TTL_SECONDS: Final = int(os.getenv("LITELLM_PARALLEL_REQUEST_SLOT_TTL_SECONDS", "3600"))
|
||||
PARALLEL_REQUEST_SLOT_TTL_SECONDS: Final = max(1, int(os.getenv("LITELLM_PARALLEL_REQUEST_SLOT_TTL_SECONDS", "3600")))
|
||||
|
||||
|
||||
CacheCounterValue: TypeAlias = int | float | str | bytes
|
||||
|
|
|
|||
|
|
@ -4960,7 +4960,7 @@ def test_pre_call_redacts_and_masks_raw_request(logging_obj):
|
|||
assert "key=*****" in raw_api_base
|
||||
|
||||
|
||||
def _streaming_logging_obj_with_callbacks(callbacks):
|
||||
def _streaming_logging_obj_with_callbacks(callbacks: list[CustomLogger]):
|
||||
from datetime import datetime
|
||||
|
||||
obj = LitellmLogging(
|
||||
|
|
|
|||
Loading…
Add table
Reference in a new issue