diff --git a/model_prices_and_context_window.json b/model_prices_and_context_window.json index 14d3c1ae846..195e2885770 100644 --- a/model_prices_and_context_window.json +++ b/model_prices_and_context_window.json @@ -9084,9 +9084,6 @@ "cache_read_input_token_cost": 5e-07, "input_cost_per_token": 5e-06, "litellm_provider": "anthropic", - "aliases": [ - "claude-opus-4-6-20250514" - ], "max_input_tokens": 1000000, "max_output_tokens": 128000, "max_tokens": 128000, diff --git a/proxy_server.log b/proxy_server.log new file mode 100644 index 00000000000..6aae2d1bcbf --- /dev/null +++ b/proxy_server.log @@ -0,0 +1,659 @@ +:128: RuntimeWarning: 'litellm.proxy.proxy_cli' found in sys.modules after import of package 'litellm.proxy', but prior to execution of 'litellm.proxy.proxy_cli'; this may result in unpredictable behaviour +17:27:32 - LiteLLM:ERROR: check_migration.py:100 - 🚨🚨🚨 prisma schema out of sync with db. Consider running these sql_commands to sync the two - ['ALTER TABLE "LiteLLM_BudgetTable" ADD COLUMN "allowed_models" TEXT[] DEFAULT ARRAY[]::TEXT[];', 'ALTER TABLE "LiteLLM_MCPServerTable" ADD COLUMN "instructions" TEXT;', 'ALTER TABLE "LiteLLM_TeamTable" ADD COLUMN "budget_limits" JSONB, ADD COLUMN "default_team_member_models" TEXT[] DEFAULT ARRAY[]::TEXT[];', 'ALTER TABLE "LiteLLM_VerificationToken" ADD COLUMN "budget_limits" JSONB;', 'CREATE INDEX "LiteLLM_HealthCheckTable_model_id_model_name_checked_at_idx" ON "LiteLLM_HealthCheckTable"("model_id", "model_name", "checked_at" DESC);'] +NoneType: None +INFO: Started server process [66039] +INFO: Waiting for application startup. +17:27:32 - LiteLLM Proxy:DEBUG: proxy_server.py:804 - litellm.proxy.proxy_server.py::startup() - CHECKING PREMIUM USER - True +17:27:32 - LiteLLM Proxy:DEBUG: proxy_server.py:817 - worker_config: {"model": null, "alias": null, "api_base": null, "api_version": "2025-02-01-preview", "debug": false, "detailed_debug": true, "temperature": null, "max_tokens": null, "request_timeout": null, "max_budget": null, "telemetry": true, "drop_params": false, "add_function_to_prompt": false, "headers": null, "save": false, "config": "proxy_server_config.yaml", "use_queue": false} +Changes to DB Schema detected +Required SQL commands: +ALTER TABLE "LiteLLM_BudgetTable" ADD COLUMN "allowed_models" TEXT[] DEFAULT ARRAY[]::TEXT[]; +ALTER TABLE "LiteLLM_MCPServerTable" ADD COLUMN "instructions" TEXT; +ALTER TABLE "LiteLLM_TeamTable" ADD COLUMN "budget_limits" JSONB, ADD COLUMN "default_team_member_models" TEXT[] DEFAULT ARRAY[]::TEXT[]; +ALTER TABLE "LiteLLM_VerificationToken" ADD COLUMN "budget_limits" JSONB; +CREATE INDEX "LiteLLM_HealthCheckTable_model_id_model_name_checked_at_idx" ON "LiteLLM_HealthCheckTable"("model_id", "model_name", "checked_at" DESC); + + β–ˆβ–ˆβ•— β–ˆβ–ˆβ•—β–ˆβ–ˆβ–ˆβ–ˆβ–ˆβ–ˆβ–ˆβ–ˆβ•—β–ˆβ–ˆβ–ˆβ–ˆβ–ˆβ–ˆβ–ˆβ•—β–ˆβ–ˆβ•— β–ˆβ–ˆβ•— β–ˆβ–ˆβ–ˆβ•— β–ˆβ–ˆβ–ˆβ•— + β–ˆβ–ˆβ•‘ β–ˆβ–ˆβ•‘β•šβ•β•β–ˆβ–ˆβ•”β•β•β•β–ˆβ–ˆβ•”β•β•β•β•β•β–ˆβ–ˆβ•‘ β–ˆβ–ˆβ•‘ β–ˆβ–ˆβ–ˆβ–ˆβ•— β–ˆβ–ˆβ–ˆβ–ˆβ•‘ + β–ˆβ–ˆβ•‘ β–ˆβ–ˆβ•‘ β–ˆβ–ˆβ•‘ β–ˆβ–ˆβ–ˆβ–ˆβ–ˆβ•— β–ˆβ–ˆβ•‘ β–ˆβ–ˆβ•‘ β–ˆβ–ˆβ•”β–ˆβ–ˆβ–ˆβ–ˆβ•”β–ˆβ–ˆβ•‘ + β–ˆβ–ˆβ•‘ β–ˆβ–ˆβ•‘ β–ˆβ–ˆβ•‘ β–ˆβ–ˆβ•”β•β•β• β–ˆβ–ˆβ•‘ β–ˆβ–ˆβ•‘ β–ˆβ–ˆβ•‘β•šβ–ˆβ–ˆβ•”β•β–ˆβ–ˆβ•‘ + β–ˆβ–ˆβ–ˆβ–ˆβ–ˆβ–ˆβ–ˆβ•—β–ˆβ–ˆβ•‘ β–ˆβ–ˆβ•‘ β–ˆβ–ˆβ–ˆβ–ˆβ–ˆβ–ˆβ–ˆβ•—β–ˆβ–ˆβ–ˆβ–ˆβ–ˆβ–ˆβ–ˆβ•—β–ˆβ–ˆβ–ˆβ–ˆβ–ˆβ–ˆβ–ˆβ•—β–ˆβ–ˆβ•‘ β•šβ•β• β–ˆβ–ˆβ•‘ + β•šβ•β•β•β•β•β•β•β•šβ•β• β•šβ•β• β•šβ•β•β•β•β•β•β•β•šβ•β•β•β•β•β•β•β•šβ•β•β•β•β•β•β•β•šβ•β• β•šβ•β• + +17:27:32 - LiteLLM:DEBUG: utils.py:5620 - checking potential_model_names in litellm.model_cost: {'split_model': '981a6c3e187a3f563db301b17f0776519148f13a3a80f13658446ed9332f8f22', 'combined_model_name': '981a6c3e187a3f563db301b17f0776519148f13a3a80f13658446ed9332f8f22', 'stripped_model_name': '981a6c3e187a3f563db301b17f0776519148f13a3a80f13658446ed9332f8f22', 'combined_stripped_model_name': '981a6c3e187a3f563db301b17f0776519148f13a3a80f13658446ed9332f8f22', 'custom_llm_provider': None} +17:27:32 - LiteLLM:DEBUG: utils.py:5915 - Error getting model info: This model isn't mapped yet. Add it here - https://github.com/BerriAI/litellm/blob/main/model_prices_and_context_window.json +17:27:32 - LiteLLM:DEBUG: utils.py:2900 - added/updated model=981a6c3e187a3f563db301b17f0776519148f13a3a80f13658446ed9332f8f22 in litellm.model_cost: 981a6c3e187a3f563db301b17f0776519148f13a3a80f13658446ed9332f8f22 +17:27:32 - LiteLLM:DEBUG: utils.py:5620 - checking potential_model_names in litellm.model_cost: {'split_model': 'oia-gpt-realtime', 'combined_model_name': 'azure/oia-gpt-realtime', 'stripped_model_name': 'azure/oia-gpt-realtime', 'combined_stripped_model_name': 'azure/oia-gpt-realtime', 'custom_llm_provider': 'azure'} +17:27:32 - LiteLLM:DEBUG: utils.py:5915 - Error getting model info: This model isn't mapped yet. Add it here - https://github.com/BerriAI/litellm/blob/main/model_prices_and_context_window.json +17:27:32 - LiteLLM:DEBUG: utils.py:2900 - added/updated model=azure/oia-gpt-realtime in litellm.model_cost: azure/oia-gpt-realtime +17:27:32 - LiteLLM:DEBUG: utils.py:5620 - checking potential_model_names in litellm.model_cost: {'split_model': '91a8591ffdee6edec7016031dc8a092b14ce59f85b1b105c0afa69280775d380', 'combined_model_name': '91a8591ffdee6edec7016031dc8a092b14ce59f85b1b105c0afa69280775d380', 'stripped_model_name': '91a8591ffdee6edec7016031dc8a092b14ce59f85b1b105c0afa69280775d380', 'combined_stripped_model_name': '91a8591ffdee6edec7016031dc8a092b14ce59f85b1b105c0afa69280775d380', 'custom_llm_provider': None} +17:27:32 - LiteLLM:DEBUG: utils.py:5915 - Error getting model info: This model isn't mapped yet. Add it here - https://github.com/BerriAI/litellm/blob/main/model_prices_and_context_window.json +17:27:32 - LiteLLM:DEBUG: utils.py:2900 - added/updated model=91a8591ffdee6edec7016031dc8a092b14ce59f85b1b105c0afa69280775d380 in litellm.model_cost: 91a8591ffdee6edec7016031dc8a092b14ce59f85b1b105c0afa69280775d380 +17:27:32 - LiteLLM:DEBUG: utils.py:5620 - checking potential_model_names in litellm.model_cost: {'split_model': 'claude-opus-4-7', 'combined_model_name': 'anthropic/claude-opus-4-7', 'stripped_model_name': 'anthropic/claude-opus-4-7', 'combined_stripped_model_name': 'anthropic/claude-opus-4-7', 'custom_llm_provider': 'anthropic'} +17:27:32 - LiteLLM:DEBUG: utils.py:2900 - added/updated model=claude-opus-4-7 in litellm.model_cost: claude-opus-4-7 +17:27:32 - LiteLLM Router:DEBUG: router.py:7008 - +Initialized Model List ['gpt-realtime', 'claude-opus-4-7'] +17:27:32 - LiteLLM Router:INFO: router.py:802 - Routing strategy: simple-shuffle +17:27:32 - LiteLLM Proxy:DEBUG: proxy_server.py:3838 - Policy engine: no policies in config, skipping +17:27:32 - LiteLLM Proxy:DEBUG: utils.py:2455 - Creating Prisma Client.. +17:27:32 - LiteLLM Proxy:DEBUG: utils.py:2522 - Success - Created Prisma Client +17:27:32 - LiteLLM Proxy:DEBUG: utils.py:3740 - PrismaClient: connect() called Attempting to Connect to DB +17:27:32 - LiteLLM Proxy:DEBUG: utils.py:3744 - PrismaClient: DB not connected, Attempting to Connect to DB +query-engine ac9d7041ed77bcc8a8dbd2ab6616b39013829574 +17:27:32 - LiteLLM Proxy:DEBUG: prisma_client.py:247 - IAM token auth not enabled, skipping token refresh task +17:27:32 - LiteLLM Proxy:INFO: utils.py:4302 - Started Prisma DB health watchdog (interval=30s, reconnect_cooldown=15s, probe_timeout=5.0s, reconnect_timeout=30.0s) +17:27:32 - LiteLLM Proxy:INFO: utils.py:4057 - Found prisma-query-engine at PID 66228. +17:27:32 - LiteLLM Proxy:INFO: utils.py:4061 - Watching engine PID 66228 via waitpid thread. +17:27:32 - LiteLLM:DEBUG: logging_callback_manager.py:336 - Custom logger of type SkillsInjectionHook, key: SkillsInjectionHook-max_iterations=10-sandbox_timeout=120-message_logging=True-turn_off_message_logging=False already exists in [, , , , , , ], not adding again.. +17:27:32 - LiteLLM Proxy:DEBUG: proxy_server.py:896 - About to initialize semantic tool filter +17:27:32 - LiteLLM Proxy:DEBUG: proxy_server.py:899 - litellm_settings keys = [] +17:27:32 - LiteLLM Proxy:DEBUG: proxy_server.py:6143 - Semantic tool filter not configured or not enabled, skipping initialization +17:27:32 - LiteLLM Proxy:DEBUG: proxy_server.py:906 - After semantic tool filter initialization +17:27:32 - LiteLLM Proxy:DEBUG: proxy_server.py:920 - prisma_client: +17:27:32 - LiteLLM Proxy:INFO: proxy_server.py:6390 - Tag spend update job scheduled at 34s interval (2.3x main job interval) +17:27:32 - LiteLLM Proxy:DEBUG: hanging_request_check.py:148 - Checking for hanging requests.... +17:27:32 - LiteLLM Proxy:INFO: utils.py:5078 - Starting spend logs queue monitor (threshold: 100, poll_interval: 2.0s) +17:27:32 - LiteLLM Proxy:INFO: utils.py:2629 - All necessary views exist! +17:27:32 - LiteLLM Proxy:INFO: proxy_server.py:875 - Password migration: No plaintext passwords found +17:27:32 - LiteLLM Proxy:DEBUG: proxy_server.py:4179 - len new_models: 0 +17:27:32 - LiteLLM Proxy:DEBUG: proxy_server.py:5384 - guardrails from the DB [] +17:27:32 - LiteLLM Proxy:INFO: policy_registry.py:577 - Synced 0 production policies and 0 draft/published (by ID) from DB to in-memory registry +17:27:32 - LiteLLM Proxy:INFO: attachment_registry.py:481 - Synced 0 attachments from DB to in-memory registry +17:27:32 - LiteLLM Proxy:DEBUG: proxy_server.py:5420 - Successfully synced policies and attachments from DB +17:27:32 - LiteLLM:DEBUG: mcp_server_manager.py:2644 - Loading MCP servers from database into registry... +17:27:32 - LiteLLM:INFO: mcp_server_manager.py:2664 - Found 0 MCP servers in database +17:27:32 - LiteLLM:DEBUG: mcp_server_manager.py:2697 - MCP registry refreshed (0 servers in registry) +17:27:32 - LiteLLM Proxy:DEBUG: pass_through_endpoints.py:2289 - initializing pass through endpoints +17:27:32 - LiteLLM Proxy:DEBUG: prompt_registry.py:54 - Discovering prompt integrations in: litellm.integrations.weave +17:27:32 - LiteLLM Proxy:DEBUG: prompt_registry.py:54 - Discovering prompt integrations in: litellm.integrations.litellm_agent +17:27:32 - LiteLLM Proxy:DEBUG: prompt_registry.py:54 - Discovering prompt integrations in: litellm.integrations.dotprompt +17:27:32 - LiteLLM Proxy:DEBUG: prompt_registry.py:65 - Found prompt_initializer_registry in litellm.integrations.dotprompt: ['dotprompt'] +17:27:32 - LiteLLM Proxy:DEBUG: prompt_registry.py:54 - Discovering prompt integrations in: litellm.integrations.gitlab +17:27:32 - LiteLLM Proxy:DEBUG: prompt_registry.py:65 - Found prompt_initializer_registry in litellm.integrations.gitlab: ['gitlab'] +17:27:32 - LiteLLM Proxy:DEBUG: prompt_registry.py:54 - Discovering prompt integrations in: litellm.integrations.azure_sentinel +17:27:32 - LiteLLM Proxy:DEBUG: prompt_registry.py:54 - Discovering prompt integrations in: litellm.integrations.arize +17:27:32 - LiteLLM Proxy:DEBUG: prompt_registry.py:65 - Found prompt_initializer_registry in litellm.integrations.arize: ['arize_phoenix'] +17:27:32 - LiteLLM Proxy:DEBUG: prompt_registry.py:54 - Discovering prompt integrations in: litellm.integrations.agentops +17:27:32 - LiteLLM Proxy:DEBUG: prompt_registry.py:54 - Discovering prompt integrations in: litellm.integrations.focus +17:27:32 - LiteLLM Proxy:DEBUG: prompt_registry.py:54 - Discovering prompt integrations in: litellm.integrations.prometheus_helpers +17:27:32 - LiteLLM Proxy:DEBUG: prompt_registry.py:54 - Discovering prompt integrations in: litellm.integrations.generic_prompt_management +17:27:32 - LiteLLM Proxy:DEBUG: prompt_registry.py:65 - Found prompt_initializer_registry in litellm.integrations.generic_prompt_management: ['generic_prompt_management'] +17:27:32 - LiteLLM Proxy:DEBUG: prompt_registry.py:54 - Discovering prompt integrations in: litellm.integrations.levo +17:27:32 - LiteLLM Proxy:DEBUG: prompt_registry.py:54 - Discovering prompt integrations in: litellm.integrations.websearch_interception +17:27:32 - LiteLLM Proxy:DEBUG: prompt_registry.py:54 - Discovering prompt integrations in: litellm.integrations.deepeval +17:27:32 - LiteLLM Proxy:DEBUG: prompt_registry.py:54 - Discovering prompt integrations in: litellm.integrations.bitbucket +17:27:32 - LiteLLM Proxy:DEBUG: prompt_registry.py:65 - Found prompt_initializer_registry in litellm.integrations.bitbucket: ['bitbucket'] +17:27:32 - LiteLLM Proxy:DEBUG: prompt_registry.py:54 - Discovering prompt integrations in: litellm.integrations.vantage +17:27:32 - LiteLLM Proxy:DEBUG: prompt_registry.py:76 - Discovered 5 prompt initializers: ['dotprompt', 'gitlab', 'arize_phoenix', 'generic_prompt_management', 'bitbucket'] +17:27:32 - LiteLLM Proxy:INFO: proxy_server.py:5563 - Loading 0 search tool(s) from database into router +17:27:32 - LiteLLM Proxy:DEBUG: proxy_server.py:5583 - No search tools found in database, keeping config-loaded search tools (if any) +17:27:32 - LiteLLM Proxy:INFO: tool_registry_writer.py:329 - ToolPolicyRegistry: synced 15 tool policies and 0 object permissions from DB +17:27:32 - LiteLLM Proxy:DEBUG: proxy_server.py:5440 - Successfully synced tool policy from DB +17:27:32 - LiteLLM:DEBUG: focus_logger.py:167 - No Focus export logger registered; skipping scheduler +17:27:32 - LiteLLM Proxy:DEBUG: proxy_server.py:6678 - key_rotation_enabled: False +17:27:32 - LiteLLM Proxy:DEBUG: proxy_server.py:6714 - Key rotation disabled (set LITELLM_KEY_ROTATION_ENABLED=true to enable) +17:27:32 - LiteLLM Proxy:INFO: proxy_server.py:6545 - Batch cost check job scheduled successfully +17:27:32 - LiteLLM Proxy:INFO: proxy_server.py:6576 - Responses cost check job scheduled successfully +17:27:32 - LiteLLM Proxy:INFO: proxy_server.py:6595 - APScheduler started with memory leak prevention settings: removed jitter, increased intervals, misfire_grace_time=3600 +17:27:32 - LiteLLM Proxy:DEBUG: proxy_server.py:6880 - LiteLLM: Pyroscope profiling is disabled (set LITELLM_ENABLE_PYROSCOPE=true to enable). +17:27:32 - LiteLLM Proxy:INFO: proxy_server.py:756 - SESSION REUSE: Created shared aiohttp session for connection pooling (ID: 4796655632, limit=1000, limit_per_host=500) +INFO: Application startup complete. +INFO: Uvicorn running on http://0.0.0.0:4000 (Press CTRL+C to quit) +17:27:39 - LiteLLM Proxy:DEBUG: http_parsing_utils.py:529 - populate_request_with_path_params: No vector_store_id present in path=/v1/chat/completions +17:27:39 - LiteLLM:DEBUG: user_api_key_auth.py:1032 - api key not found in cache. +17:27:39 - LiteLLM Proxy:DEBUG: common_request_processing.py:876 - Request received by LiteLLM: +{ + "model": "claude-opus-4-7", + "messages": [ + { + "role": "user", + "content": "Hello, how are you?" + } + ], + "reasoning_effort": "max", + "metadata": { + "user_api_key_user_id": "default_user_id" + } +} +17:27:39 - LiteLLM Proxy:DEBUG: litellm_pre_call_utils.py:971 - Request Headers: {'content-type': 'application/json', 'user-agent': 'PostmanRuntime/7.53.0', 'accept': '*/*', 'postman-token': '08d81a08-addf-4a27-a9f8-302f317d7296', 'host': '0.0.0.0:4000', 'accept-encoding': 'gzip, deflate, br', 'connection': 'keep-alive', 'content-length': '183'} +17:27:39 - LiteLLM Proxy:DEBUG: litellm_pre_call_utils.py:972 - Raw Headers: RedactedDict(REDACTED) +17:27:39 - LiteLLM Proxy:DEBUG: litellm_pre_call_utils.py:1065 - receiving data: {'model': 'claude-opus-4-7', 'messages': [{'role': 'user', 'content': 'Hello, how are you?'}], 'reasoning_effort': 'max', 'metadata': {'headers': {'content-type': 'application/json', 'user-agent': 'PostmanRuntime/7.53.0', 'accept': '*/*', 'postman-token': '08d81a08-addf-4a27-a9f8-302f317d7296', 'host': '0.0.0.0:4000', 'accept-encoding': 'gzip, deflate, br', 'connection': 'keep-alive', 'content-length': '183'}}, 'proxy_server_request': {'url': 'http://0.0.0.0:4000/v1/chat/completions', 'method': 'POST', 'headers': {'content-type': 'application/json', 'user-agent': 'PostmanRuntime/7.53.0', 'accept': '*/*', 'postman-token': '08d81a08-addf-4a27-a9f8-302f317d7296', 'host': '0.0.0.0:4000', 'accept-encoding': 'gzip, deflate, br', 'connection': 'keep-alive', 'content-length': '183'}, 'body': None, 'arrival_time': 1776686259.683504}, 'secret_fields': {'raw_headers': RedactedDict(REDACTED)}} +17:27:39 - LiteLLM Proxy:DEBUG: litellm_pre_call_utils.py:2246 - Policy engine: registry initialized=True, policy_count=0 +17:27:39 - LiteLLM Proxy:DEBUG: litellm_pre_call_utils.py:2267 - Policy engine: matching policies for context team_alias=None, key_alias=None, model=claude-opus-4-7, tags=None +17:27:39 - LiteLLM Proxy:DEBUG: litellm_pre_call_utils.py:2104 - Policy engine: matched policies via attachments: [] +17:27:39 - LiteLLM Proxy:DEBUG: policy_resolver.py:173 - No policies match context: team_alias=None, key_alias=None, model=claude-opus-4-7 +17:27:39 - LiteLLM Proxy:DEBUG: litellm_pre_call_utils.py:2165 - Policy engine: resolved guardrails: [] +17:27:39 - LiteLLM Proxy:DEBUG: litellm_pre_call_utils.py:1398 - [PROXY] returned data from litellm_pre_call_utils: {'model': 'claude-opus-4-7', 'messages': [{'role': 'user', 'content': 'Hello, how are you?'}], 'reasoning_effort': 'max', 'metadata': {'headers': {'content-type': 'application/json', 'user-agent': 'PostmanRuntime/7.53.0', 'accept': '*/*', 'postman-token': '08d81a08-addf-4a27-a9f8-302f317d7296', 'host': '0.0.0.0:4000', 'accept-encoding': 'gzip, deflate, br', 'connection': 'keep-alive', 'content-length': '183'}, 'requester_metadata': {'headers': {'content-type': 'application/json', 'user-agent': 'PostmanRuntime/7.53.0', 'accept': '*/*', 'postman-token': '08d81a08-addf-4a27-a9f8-302f317d7296', 'host': '0.0.0.0:4000', 'accept-encoding': 'gzip, deflate, br', 'connection': 'keep-alive', 'content-length': '183'}}, 'user_api_key_hash': '88dc28d0f030c55ed4ab77ed8faf098196cb1c05df778539800c9f1243fe6b4b', 'user_api_key_alias': None, 'user_api_key_spend': 0.0, 'user_api_key_max_budget': None, 'user_api_key_team_id': None, 'user_api_key_project_id': None, 'user_api_key_project_alias': None, 'user_api_key_user_id': 'default_user_id', 'user_api_key_org_id': None, 'user_api_key_org_alias': None, 'user_api_key_team_alias': None, 'user_api_key_end_user_id': None, 'user_api_key_user_email': None, 'user_api_key_request_route': '/v1/chat/completions', 'user_api_key_budget_reset_at': None, 'user_api_key_auth_metadata': {}, 'user_REDACTED', 'agent_id': None, 'user_api_end_user_max_budget': None, 'user_api_key_auth': UserAPIKeyAuth(token='88dc28d0f030c55ed4ab77ed8faf098196cb1c05df778539800c9f1243fe6b4b', key_name=None, key_alias=None, spend=0.0, max_budget=None, expires=None, models=[], aliases={}, config={}, user_id='default_user_id', team_id=None, agent_id=None, project_id=None, max_parallel_requests=None, metadata={}, tpm_limit=None, rpm_limit=None, budget_duration=None, budget_reset_at=None, allowed_cache_controls=[], allowed_routes=[], permissions={}, model_spend={}, model_max_budget={}, soft_budget_cooldown=False, blocked=None, litellm_budget_table=None, org_id=None, created_at=None, created_by=None, updated_at=None, updated_by=None, last_active=None, object_permission_id=None, object_permission=None, access_group_ids=None, rotation_count=0, auto_rotate=False, rotation_interval=None, last_rotation_at=None, key_rotation_at=None, router_settings=None, budget_limits=None, team_spend=None, team_alias=None, team_tpm_limit=None, team_rpm_limit=None, team_max_budget=None, team_soft_budget=None, team_models=[], team_blocked=False, soft_budget=None, team_model_aliases=None, team_member=None, team_metadata=None, team_object_permission_id=None, team_member_spend=None, team_member_tpm_limit=None, team_member_rpm_limit=None, end_user_id=None, end_user_tpm_limit=None, end_user_rpm_limit=None, end_user_max_budget=None, end_user_model_max_budget=None, organization_alias=None, organization_max_budget=None, organization_tpm_limit=None, organization_rpm_limit=None, organization_metadata=None, project_alias=None, project_metadata=None, last_refreshed_at=None, REDACTED', user_role=, allowed_model_region=None, parent_otel_span=None, rpm_limit_per_model=None, tpm_limit_per_model=None, user_tpm_limit=None, user_rpm_limit=None, user_email=None, user_spend=None, user_max_budget=None, request_route='/v1/chat/completions', user=None, created_by_user=None, end_user_object_permission=None, jwt_claims=None), 'litellm_api_version': '1.83.9', 'global_max_parallel_requests': None, 'user_api_key_team_max_budget': None, 'user_api_key_team_spend': None, 'user_api_key_model_max_budget': {}, 'user_api_key_end_user_model_max_budget': None, 'user_api_key_user_spend': None, 'user_api_key_user_max_budget': None, 'user_api_key_metadata': {}, 'user_api_key_team_metadata': None, 'user_api_key_object_permission_id': None, 'user_api_key_team_object_permission_id': None, 'endpoint': 'http://0.0.0.0:4000/v1/chat/completions', 'litellm_parent_otel_span': None, 'requester_ip_address': '127.0.0.1', 'user_agent': 'PostmanRuntime/7.53.0'}, 'proxy_server_request': {'url': 'http://0.0.0.0:4000/v1/chat/completions', 'method': 'POST', 'headers': {'content-type': 'application/json', 'user-agent': 'PostmanRuntime/7.53.0', 'accept': '*/*', 'postman-token': '08d81a08-addf-4a27-a9f8-302f317d7296', 'host': '0.0.0.0:4000', 'accept-encoding': 'gzip, deflate, br', 'connection': 'keep-alive', 'content-length': '183'}, 'body': {'model': 'claude-opus-4-7', 'messages': [{'role': 'user', 'content': 'Hello, how are you?'}], 'reasoning_effort': 'max', 'metadata': {'headers': {'content-type': 'application/json', 'user-agent': 'PostmanRuntime/7.53.0', 'accept': '*/*', 'postman-token': '08d81a08-addf-4a27-a9f8-302f317d7296', 'host': '0.0.0.0:4000', 'accept-encoding': 'gzip, deflate, br', 'connection': 'keep-alive', 'content-length': '183'}, 'requester_metadata': {'headers': {'content-type': 'application/json', 'user-agent': 'PostmanRuntime/7.53.0', 'accept': '*/*', 'postman-token': '08d81a08-addf-4a27-a9f8-302f317d7296', 'host': '0.0.0.0:4000', 'accept-encoding': 'gzip, deflate, br', 'connection': 'keep-alive', 'content-length': '183'}}, 'user_api_key_hash': '88dc28d0f030c55ed4ab77ed8faf098196cb1c05df778539800c9f1243fe6b4b', 'user_api_key_alias': None, 'user_api_key_spend': 0.0, 'user_api_key_max_budget': None, 'user_api_key_team_id': None, 'user_api_key_project_id': None, 'user_api_key_project_alias': None, 'user_api_key_user_id': 'default_user_id', 'user_api_key_org_id': None, 'user_api_key_org_alias': None, 'user_api_key_team_alias': None, 'user_api_key_end_user_id': None, 'user_api_key_user_email': None, 'user_api_key_request_route': '/v1/chat/completions', 'user_api_key_budget_reset_at': None, 'user_api_key_auth_metadata': {}, 'user_REDACTED', 'agent_id': None, 'user_api_end_user_max_budget': None, 'user_api_key_auth': UserAPIKeyAuth(token='88dc28d0f030c55ed4ab77ed8faf098196cb1c05df778539800c9f1243fe6b4b', key_name=None, key_alias=None, spend=0.0, max_budget=None, expires=None, models=[], aliases={}, config={}, user_id='default_user_id', team_id=None, agent_id=None, project_id=None, max_parallel_requests=None, metadata={}, tpm_limit=None, rpm_limit=None, budget_duration=None, budget_reset_at=None, allowed_cache_controls=[], allowed_routes=[], permissions={}, model_spend={}, model_max_budget={}, soft_budget_cooldown=False, blocked=None, litellm_budget_table=None, org_id=None, created_at=None, created_by=None, updated_at=None, updated_by=None, last_active=None, object_permission_id=None, object_permission=None, access_group_ids=None, rotation_count=0, auto_rotate=False, rotation_interval=None, last_rotation_at=None, key_rotation_at=None, router_settings=None, budget_limits=None, team_spend=None, team_alias=None, team_tpm_limit=None, team_rpm_limit=None, team_max_budget=None, team_soft_budget=None, team_models=[], team_blocked=False, soft_budget=None, team_model_aliases=None, team_member=None, team_metadata=None, team_object_permission_id=None, team_member_spend=None, team_member_tpm_limit=None, team_member_rpm_limit=None, end_user_id=None, end_user_tpm_limit=None, end_user_rpm_limit=None, end_user_max_budget=None, end_user_model_max_budget=None, organization_alias=None, organization_max_budget=None, organization_tpm_limit=None, organization_rpm_limit=None, organization_metadata=None, project_alias=None, project_metadata=None, last_refreshed_at=None, REDACTED', user_role=, allowed_model_region=None, parent_otel_span=None, rpm_limit_per_model=None, tpm_limit_per_model=None, user_tpm_limit=None, user_rpm_limit=None, user_email=None, user_spend=None, user_max_budget=None, request_route='/v1/chat/completions', user=None, created_by_user=None, end_user_object_permission=None, jwt_claims=None), 'litellm_api_version': '1.83.9', 'global_max_parallel_requests': None, 'user_api_key_team_max_budget': None, 'user_api_key_team_spend': None, 'user_api_key_model_max_budget': {}, 'user_api_key_end_user_model_max_budget': None, 'user_api_key_user_spend': None, 'user_api_key_user_max_budget': None, 'user_api_key_metadata': {}, 'user_api_key_team_metadata': None, 'user_api_key_object_permission_id': None, 'user_api_key_team_object_permission_id': None, 'endpoint': 'http://0.0.0.0:4000/v1/chat/completions', 'litellm_parent_otel_span': None, 'requester_ip_address': '127.0.0.1', 'user_agent': 'PostmanRuntime/7.53.0'}, 'proxy_server_request': {...}, 'secret_fields': {'raw_headers': RedactedDict(REDACTED)}}, 'arrival_time': 1776686259.683504}, 'secret_fields': {'raw_headers': RedactedDict(REDACTED)}} +17:27:39 - LiteLLM:DEBUG: logging_callback_manager.py:336 - Custom logger of type _ProxyDBLogger, key: _ProxyDBLogger-message_logging=True-turn_off_message_logging=False already exists in [>, , , ], not adding again.. +17:27:39 - LiteLLM:DEBUG: utils.py:482 - Initialized litellm callbacks, Async Success Callbacks: [>, , , , , , , , , , , , ] +17:27:39 - LiteLLM:DEBUG: utils.py:1008 - Removing thought signatures from tool call IDs for non-Gemini model +17:27:39 - LiteLLM:DEBUG: litellm_logging.py:557 - self.optional_params: {} +17:27:39 - LiteLLM Proxy:DEBUG: utils.py:1346 - Inside Proxy Logging Pre-call hook! +17:27:39 - LiteLLM Proxy:DEBUG: max_budget_limiter.py:23 - Inside Max Budget Limiter Pre-Call Hook +17:27:39 - LiteLLM Proxy:DEBUG: parallel_request_limiter_v3.py:1288 - Inside Rate Limit Pre-Call Hook +17:27:39 - LiteLLM Proxy:DEBUG: cache_control_check.py:27 - Inside Cache Control Check Pre-Call Hook +17:27:39 - LiteLLM Proxy:INFO: route_llm_request.py:164 - SESSION REUSE: Attached shared aiohttp session to request (ID: 4796655632) +17:27:39 - LiteLLM Router:DEBUG: router.py:5679 - Inside async function with retries. +17:27:39 - LiteLLM Router:DEBUG: router.py:5703 - async function w/ retries: original_function - >, num_retries - 2 +17:27:39 - LiteLLM Router:DEBUG: router.py:9140 - initial list of deployments: [{'model_name': 'claude-opus-4-7', 'litellm_params': {'use_in_pass_through': False, 'use_litellm_proxy': False, 'merge_reasoning_content_in_choices': False, 'model': 'anthropic/claude-opus-4-7'}, 'model_info': {'id': '91a8591ffdee6edec7016031dc8a092b14ce59f85b1b105c0afa69280775d380', 'db_model': False}}] +17:27:39 - LiteLLM Router:DEBUG: router.py:9216 - healthy_deployments after team filter: [{'model_name': 'claude-opus-4-7', 'litellm_params': {'use_in_pass_through': False, 'use_litellm_proxy': False, 'merge_reasoning_content_in_choices': False, 'model': 'anthropic/claude-opus-4-7'}, 'model_info': {'id': '91a8591ffdee6edec7016031dc8a092b14ce59f85b1b105c0afa69280775d380', 'db_model': False}}] +17:27:39 - LiteLLM Router:DEBUG: router.py:9226 - healthy_deployments after web search filter: [{'model_name': 'claude-opus-4-7', 'litellm_params': {'use_in_pass_through': False, 'use_litellm_proxy': False, 'merge_reasoning_content_in_choices': False, 'model': 'anthropic/claude-opus-4-7'}, 'model_info': {'id': '91a8591ffdee6edec7016031dc8a092b14ce59f85b1b105c0afa69280775d380', 'db_model': False}}] +17:27:39 - LiteLLM Router:DEBUG: cooldown_handlers.py:347 - retrieve cooldown models: [] +17:27:39 - LiteLLM Router:DEBUG: router.py:9245 - cooldown deployments: [] +17:27:39 - LiteLLM Router:DEBUG: router.py:9987 - cooldown deployments: [] +17:27:39 - LiteLLM:DEBUG: utils.py:482 - + +17:27:39 - LiteLLM:DEBUG: utils.py:482 - Request to litellm: +17:27:39 - LiteLLM:DEBUG: utils.py:482 - litellm.acompletion(use_in_pass_through=False, use_litellm_proxy=False, merge_reasoning_content_in_choices=False, model='anthropic/claude-opus-4-7', messages=[{'role': 'user', 'content': 'Hello, how are you?'}], caching=False, client=None, reasoning_effort='max', metadata={'headers': {'content-type': 'application/json', 'user-agent': 'PostmanRuntime/7.53.0', 'accept': '*/*', 'postman-token': '08d81a08-addf-4a27-a9f8-302f317d7296', 'host': '0.0.0.0:4000', 'accept-encoding': 'gzip, deflate, br', 'connection': 'keep-alive', 'content-length': '183'}, 'requester_metadata': {'headers': {'content-type': 'application/json', 'user-agent': 'PostmanRuntime/7.53.0', 'accept': '*/*', 'postman-token': '08d81a08-addf-4a27-a9f8-302f317d7296', 'host': '0.0.0.0:4000', 'accept-encoding': 'gzip, deflate, br', 'connection': 'keep-alive', 'content-length': '183'}}, 'user_api_key_hash': '88dc28d0f030c55ed4ab77ed8faf098196cb1c05df778539800c9f1243fe6b4b', 'user_api_key_alias': None, 'user_api_key_spend': 0.0, 'user_api_key_max_budget': None, 'user_api_key_team_id': None, 'user_api_key_project_id': None, 'user_api_key_project_alias': None, 'user_api_key_user_id': 'default_user_id', 'user_api_key_org_id': None, 'user_api_key_org_alias': None, 'user_api_key_team_alias': None, 'user_api_key_end_user_id': None, 'user_api_key_user_email': None, 'user_api_key_request_route': '/v1/chat/completions', 'user_api_key_budget_reset_at': None, 'user_api_key_auth_metadata': {}, 'user_REDACTED', 'agent_id': None, 'user_api_end_user_max_budget': None, 'user_api_key_auth': UserAPIKeyAuth(token='88dc28d0f030c55ed4ab77ed8faf098196cb1c05df778539800c9f1243fe6b4b', key_name=None, key_alias=None, spend=0.0, max_budget=None, expires=None, models=[], aliases={}, config={}, user_id='default_user_id', team_id=None, agent_id=None, project_id=None, max_parallel_requests=None, metadata={}, tpm_limit=None, rpm_limit=None, budget_duration=None, budget_reset_at=None, allowed_cache_controls=[], allowed_routes=[], permissions={}, model_spend={}, model_max_budget={}, soft_budget_cooldown=False, blocked=None, litellm_budget_table=None, org_id=None, created_at=None, created_by=None, updated_at=None, updated_by=None, last_active=None, object_permission_id=None, object_permission=None, access_group_ids=None, rotation_count=0, auto_rotate=False, rotation_interval=None, last_rotation_at=None, key_rotation_at=None, router_settings=None, budget_limits=None, team_spend=None, team_alias=None, team_tpm_limit=None, team_rpm_limit=None, team_max_budget=None, team_soft_budget=None, team_models=[], team_blocked=False, soft_budget=None, team_model_aliases=None, team_member=None, team_metadata=None, team_object_permission_id=None, team_member_spend=None, team_member_tpm_limit=None, team_member_rpm_limit=None, end_user_id=None, end_user_tpm_limit=None, end_user_rpm_limit=None, end_user_max_budget=None, end_user_model_max_budget=None, organization_alias=None, organization_max_budget=None, organization_tpm_limit=None, organization_rpm_limit=None, organization_metadata=None, project_alias=None, project_metadata=None, last_refreshed_at=1776686259.698418, REDACTED', user_role=, allowed_model_region=None, parent_otel_span=None, rpm_limit_per_model=None, tpm_limit_per_model=None, user_tpm_limit=None, user_rpm_limit=None, user_email=None, user_spend=None, user_max_budget=None, request_route='/v1/chat/completions', user=None, created_by_user=None, end_user_object_permission=None, jwt_claims=None), 'litellm_api_version': '1.83.9', 'global_max_parallel_requests': None, 'user_api_key_team_max_budget': None, 'user_api_key_team_spend': None, 'user_api_key_model_max_budget': {}, 'user_api_key_end_user_model_max_budget': None, 'user_api_key_user_spend': None, 'user_api_key_user_max_budget': None, 'user_api_key_metadata': {}, 'user_api_key_team_metadata': None, 'user_api_key_object_permission_id': None, 'user_api_key_team_object_permission_id': None, 'endpoint': 'http://0.0.0.0:4000/v1/chat/completions', 'litellm_parent_otel_span': None, 'requester_ip_address': '127.0.0.1', 'user_agent': 'PostmanRuntime/7.53.0', 'queue_time_seconds': 0.009355783462524414, 'model_group': 'claude-opus-4-7', 'model_group_alias': None, 'model_group_size': 1, 'attempted_retries': 0, 'max_retries': 2, 'deployment': 'anthropic/claude-opus-4-7', 'model_info': {'id': '91a8591ffdee6edec7016031dc8a092b14ce59f85b1b105c0afa69280775d380', 'db_model': False}, 'api_base': None, 'deployment_model_name': 'claude-opus-4-7', 'caching_groups': None}, proxy_server_request={'url': 'http://0.0.0.0:4000/v1/chat/completions', 'method': 'POST', 'headers': {'content-type': 'application/json', 'user-agent': 'PostmanRuntime/7.53.0', 'accept': '*/*', 'postman-token': '08d81a08-addf-4a27-a9f8-302f317d7296', 'host': '0.0.0.0:4000', 'accept-encoding': 'gzip, deflate, br', 'connection': 'keep-alive', 'content-length': '183'}, 'body': {'model': 'claude-opus-4-7', 'messages': [{'role': 'user', 'content': 'Hello, how are you?'}], 'reasoning_effort': 'max', 'metadata': {'headers': {'content-type': 'application/json', 'user-agent': 'PostmanRuntime/7.53.0', 'accept': '*/*', 'postman-token': '08d81a08-addf-4a27-a9f8-302f317d7296', 'host': '0.0.0.0:4000', 'accept-encoding': 'gzip, deflate, br', 'connection': 'keep-alive', 'content-length': '183'}, 'requester_metadata': {'headers': {'content-type': 'application/json', 'user-agent': 'PostmanRuntime/7.53.0', 'accept': '*/*', 'postman-token': '08d81a08-addf-4a27-a9f8-302f317d7296', 'host': '0.0.0.0:4000', 'accept-encoding': 'gzip, deflate, br', 'connection': 'keep-alive', 'content-length': '183'}}, 'user_api_key_hash': '88dc28d0f030c55ed4ab77ed8faf098196cb1c05df778539800c9f1243fe6b4b', 'user_api_key_alias': None, 'user_api_key_spend': 0.0, 'user_api_key_max_budget': None, 'user_api_key_team_id': None, 'user_api_key_project_id': None, 'user_api_key_project_alias': None, 'user_api_key_user_id': 'default_user_id', 'user_api_key_org_id': None, 'user_api_key_org_alias': None, 'user_api_key_team_alias': None, 'user_api_key_end_user_id': None, 'user_api_key_user_email': None, 'user_api_key_request_route': '/v1/chat/completions', 'user_api_key_budget_reset_at': None, 'user_api_key_auth_metadata': {}, 'user_REDACTED', 'agent_id': None, 'user_api_end_user_max_budget': None, 'user_api_key_auth': UserAPIKeyAuth(token='88dc28d0f030c55ed4ab77ed8faf098196cb1c05df778539800c9f1243fe6b4b', key_name=None, key_alias=None, spend=0.0, max_budget=None, expires=None, models=[], aliases={}, config={}, user_id='default_user_id', team_id=None, agent_id=None, project_id=None, max_parallel_requests=None, metadata={}, tpm_limit=None, rpm_limit=None, budget_duration=None, budget_reset_at=None, allowed_cache_controls=[], allowed_routes=[], permissions={}, model_spend={}, model_max_budget={}, soft_budget_cooldown=False, blocked=None, litellm_budget_table=None, org_id=None, created_at=None, created_by=None, updated_at=None, updated_by=None, last_active=None, object_permission_id=None, object_permission=None, access_group_ids=None, rotation_count=0, auto_rotate=False, rotation_interval=None, last_rotation_at=None, key_rotation_at=None, router_settings=None, budget_limits=None, team_spend=None, team_alias=None, team_tpm_limit=None, team_rpm_limit=None, team_max_budget=None, team_soft_budget=None, team_models=[], team_blocked=False, soft_budget=None, team_model_aliases=None, team_member=None, team_metadata=None, team_object_permission_id=None, team_member_spend=None, team_member_tpm_limit=None, team_member_rpm_limit=None, end_user_id=None, end_user_tpm_limit=None, end_user_rpm_limit=None, end_user_max_budget=None, end_user_model_max_budget=None, organization_alias=None, organization_max_budget=None, organization_tpm_limit=None, organization_rpm_limit=None, organization_metadata=None, project_alias=None, project_metadata=None, last_refreshed_at=1776686259.698418, REDACTED', user_role=, allowed_model_region=None, parent_otel_span=None, rpm_limit_per_model=None, tpm_limit_per_model=None, user_tpm_limit=None, user_rpm_limit=None, user_email=None, user_spend=None, user_max_budget=None, request_route='/v1/chat/completions', user=None, created_by_user=None, end_user_object_permission=None, jwt_claims=None), 'litellm_api_version': '1.83.9', 'global_max_parallel_requests': None, 'user_api_key_team_max_budget': None, 'user_api_key_team_spend': None, 'user_api_key_model_max_budget': {}, 'user_api_key_end_user_model_max_budget': None, 'user_api_key_user_spend': None, 'user_api_key_user_max_budget': None, 'user_api_key_metadata': {}, 'user_api_key_team_metadata': None, 'user_api_key_object_permission_id': None, 'user_api_key_team_object_permission_id': None, 'endpoint': 'http://0.0.0.0:4000/v1/chat/completions', 'litellm_parent_otel_span': None, 'requester_ip_address': '127.0.0.1', 'user_agent': 'PostmanRuntime/7.53.0', 'queue_time_seconds': 0.009355783462524414, 'model_group': 'claude-opus-4-7', 'model_group_alias': None, 'model_group_size': 1, 'attempted_retries': 0, 'max_retries': 2, 'deployment': 'anthropic/claude-opus-4-7', 'model_info': {'id': '91a8591ffdee6edec7016031dc8a092b14ce59f85b1b105c0afa69280775d380', 'db_model': False}, 'api_base': None, 'deployment_model_name': 'claude-opus-4-7', 'caching_groups': None}, 'proxy_server_request': {...}, 'secret_fields': {'raw_headers': RedactedDict(REDACTED)}}, 'arrival_time': 1776686259.683504}, secret_fields={'raw_headers': RedactedDict(REDACTED)}, litellm_call_id='9d7c129a-8dcc-4fb1-8a30-4d7117085a4f', litellm_logging_obj=, shared_session=, stream=False, litellm_trace_id='066d410c-af0d-4e32-aca5-e3e657fbaff2', model_info={'id': '91a8591ffdee6edec7016031dc8a092b14ce59f85b1b105c0afa69280775d380', 'db_model': False}, timeout=6000.0, max_retries=0) +17:27:39 - LiteLLM:DEBUG: utils.py:482 - + +17:27:39 - LiteLLM:DEBUG: utils.py:482 - ASYNC kwargs[caching]: False; litellm.cache: None; kwargs.get('cache'): None +17:27:39 - LiteLLM:DEBUG: main.py:523 - πŸ”„ SHARED SESSION: acompletion called with shared_session (ID: 4796655632) +17:27:39 - LiteLLM:DEBUG: utils.py:5620 - checking potential_model_names in litellm.model_cost: {'split_model': 'claude-opus-4-7', 'combined_model_name': 'anthropic/claude-opus-4-7', 'stripped_model_name': 'claude-opus-4-7', 'combined_stripped_model_name': 'anthropic/claude-opus-4-7', 'custom_llm_provider': 'anthropic'} +17:27:39 - LiteLLM:INFO: utils.py:4011 - +LiteLLM completion() model= claude-opus-4-7; provider = anthropic +17:27:39 - LiteLLM:DEBUG: utils.py:4014 - +LiteLLM: Params passed to completion() {'model': 'claude-opus-4-7', 'functions': None, 'function_call': None, 'temperature': None, 'top_p': None, 'n': None, 'stream': False, 'stream_options': None, 'stop': None, 'max_tokens': None, 'max_completion_tokens': None, 'modalities': None, 'prediction': None, 'audio': None, 'presence_penalty': None, 'frequency_penalty': None, 'logit_bias': None, 'user': None, 'custom_llm_provider': 'anthropic', 'response_format': None, 'seed': None, 'tools': None, 'tool_choice': None, 'max_retries': 0, 'logprobs': None, 'top_logprobs': None, 'extra_headers': None, 'api_version': None, 'parallel_tool_calls': None, 'drop_params': None, 'allowed_openai_params': None, 'reasoning_effort': 'max', 'verbosity': None, 'additional_drop_params': None, 'messages': [{'role': 'user', 'content': 'Hello, how are you?'}], 'thinking': None, 'web_search_options': None, 'safety_identifier': None, 'service_tier': None} +17:27:39 - LiteLLM:DEBUG: utils.py:4017 - +LiteLLM: Non-Default params passed to completion() {'stream': False, 'max_retries': 0, 'reasoning_effort': 'max'} +17:27:39 - LiteLLM:DEBUG: utils.py:482 - Final returned optional params: {'thinking': {'type': 'adaptive'}, 'output_config': {'effort': 'max'}} +17:27:39 - LiteLLM:DEBUG: litellm_logging.py:557 - self.optional_params: {'stream': False, 'max_retries': 0, 'reasoning_effort': 'max'} +17:27:39 - LiteLLM:DEBUG: utils.py:5620 - checking potential_model_names in litellm.model_cost: {'split_model': 'claude-opus-4-7', 'combined_model_name': 'anthropic/claude-opus-4-7', 'stripped_model_name': 'claude-opus-4-7', 'combined_stripped_model_name': 'anthropic/claude-opus-4-7', 'custom_llm_provider': 'anthropic'} +17:27:39 - LiteLLM:DEBUG: utils.py:5620 - checking potential_model_names in litellm.model_cost: {'split_model': 'claude-opus-4-7', 'combined_model_name': 'anthropic/claude-opus-4-7', 'stripped_model_name': 'claude-opus-4-7', 'combined_stripped_model_name': 'anthropic/claude-opus-4-7', 'custom_llm_provider': 'anthropic'} +17:27:39 - LiteLLM:DEBUG: litellm_logging.py:1160 -  + +POST Request Sent from LiteLLM: +curl -X POST \ +https://api.anthropic.com/v1/messages \ +-H 'anthropic-version: 2023-06-01' -H 'accept: application/json' -H 'content-type: application/json' -H 'REDACTED' \ +-d '{'model': 'claude-opus-4-7', 'messages': [{'role': 'user', 'content': [{'type': 'text', 'text': 'Hello, how are you?'}]}], 'thinking': {'type': 'adaptive'}, 'output_config': {'effort': 'max'}, 'max_tokens': 128000}' + + +17:27:39 - LiteLLM:DEBUG: main.py:7275 - _is_function_call: False +17:27:39 - LiteLLM:DEBUG: litellm_logging.py:1233 - RAW RESPONSE: + + + +17:27:39 - LiteLLM:DEBUG: http_handler.py:840 - Using AiohttpTransport... +17:27:39 - LiteLLM:DEBUG: http_handler.py:908 - Creating AiohttpTransport... +17:27:39 - LiteLLM:DEBUG: http_handler.py:922 - NEW SESSION: Creating new ClientSession (no shared session provided) +17:27:42 - LiteLLM:DEBUG: litellm_logging.py:1233 - RAW RESPONSE: +{"model":"claude-opus-4-7","id":"msg_018BKbY3htb4ZV1AMbYuQiEe","type":"message","role":"assistant","content":[{"type":"text","text":"Hello! I'm doing well, thanks for asking. I'm ready to help with whatever you needβ€”whether that's answering questions, brainstorming ideas, working through a problem, or just having a conversation. What can I do for you today?"}],"stop_reason":"end_turn","stop_sequence":null,"stop_details":null,"usage":{"input_tokens":19,"cache_creation_input_tokens":0,"cache_read_input_tokens":0,"cache_creation":{"ephemeral_5m_input_tokens":0,"ephemeral_1h_input_tokens":0},"output_tokens":76,"service_tier":"standard","inference_geo":"global"}} + + +17:27:42 - LiteLLM:DEBUG: litellm_logging.py:3312 - Filtered callbacks: [] +17:27:42 - LiteLLM:DEBUG: cost_calculator.py:1134 - selected model name for cost calculation: anthropic/claude-opus-4-7 +17:27:42 - LiteLLM:DEBUG: utils.py:5620 - checking potential_model_names in litellm.model_cost: {'split_model': 'claude-opus-4-7', 'combined_model_name': 'anthropic/claude-opus-4-7', 'stripped_model_name': 'claude-opus-4-7', 'combined_stripped_model_name': 'anthropic/claude-opus-4-7', 'custom_llm_provider': 'anthropic'} +17:27:42 - LiteLLM:DEBUG: litellm_logging.py:1556 - response_cost: 0.001995 +17:27:42 - LiteLLM Router:INFO: router.py:2202 - litellm.acompletion(model=anthropic/claude-opus-4-7) 200 OK +17:27:42 - LiteLLM Router:DEBUG: router.py:5612 - Async Response: ModelResponse(id='chatcmpl-db0b1b06-c051-4281-98ff-6add4a98effb', created=1776686262, model='claude-opus-4-7', object='chat.completion', system_fingerprint=None, choices=[Choices(finish_reason='stop', index=0, message=Message(content="Hello! I'm doing well, thanks for asking. I'm ready to help with whatever you needβ€”whether that's answering questions, brainstorming ideas, working through a problem, or just having a conversation. What can I do for you today?", role='assistant', tool_calls=None, function_call=None, provider_specific_fields={'citations': None, 'thinking_blocks': None}))], usage=Usage(completion_tokens=76, prompt_tokens=19, total_tokens=95, completion_tokens_details=CompletionTokensDetailsWrapper(accepted_prediction_tokens=None, audio_tokens=None, reasoning_tokens=0, rejected_prediction_tokens=None, text_tokens=76, image_tokens=None, video_tokens=None), prompt_tokens_details=PromptTokensDetailsWrapper(audio_tokens=None, cached_tokens=0, text_tokens=19, image_tokens=None, video_tokens=None, cache_creation_tokens=0, cache_creation_token_details=CacheCreationTokenDetails(ephemeral_5m_input_tokens=0, ephemeral_1h_input_tokens=0)), cache_creation_input_tokens=0, cache_read_input_tokens=0, inference_geo='global', speed=None)) +17:27:42 - LiteLLM:DEBUG: utils.py:482 - Async Wrapper: Completed Call, calling async_success_handler: > +17:27:42 - LiteLLM:DEBUG: litellm_logging.py:3312 - Filtered callbacks: [] + +#------------------------------------------------------------# +# # +# 'I get frustrated when the product...' # +# https://github.com/BerriAI/litellm/issues/new # +# # +#------------------------------------------------------------# + + Thank you for using LiteLLM! - Krrish & Ishaan + + + +Give Feedback / Get Help: https://github.com/BerriAI/litellm/issues/new + + +LiteLLM: Proxy initialized with Config, Set models: + gpt-realtime + claude-opus-4-7 +INFO: 127.0.0.1:51006 - "POST /v1/chat/completions HTTP/1.1" 200 OK +17:27:42 - LiteLLM:DEBUG: utils.py:482 - Logging Details LiteLLM-Async Success Call, cache_hit=None +17:27:42 - LiteLLM:DEBUG: utils.py:5620 - checking potential_model_names in litellm.model_cost: {'split_model': 'claude-opus-4-7', 'combined_model_name': 'anthropic/claude-opus-4-7', 'stripped_model_name': 'claude-opus-4-7', 'combined_stripped_model_name': 'anthropic/claude-opus-4-7', 'custom_llm_provider': 'anthropic'} +{ + "id": "chatcmpl-db0b1b06-c051-4281-98ff-6add4a98effb", + "trace_id": "066d410c-af0d-4e32-aca5-e3e657fbaff2", + "call_type": "acompletion", + "cache_hit": null, + "stream": null, + "status": "success", + "status_fields": { + "llm_api_status": "success", + "guardrail_status": "not_run" + }, + "custom_llm_provider": "anthropic", + "saved_cache_cost": 0.0, + "startTime": 1776686259.704125, + "endTime": 1776686262.173925, + "completionStartTime": 1776686262.173925, + "response_time": 2.4697999954223633, + "model": "anthropic/claude-opus-4-7", + "metadata": { + "user_api_key_hash": "88dc28d0f030c55ed4ab77ed8faf098196cb1c05df778539800c9f1243fe6b4b", + "user_api_key_alias": null, + "user_api_key_spend": 0.0, + "user_api_key_max_budget": null, + "user_api_key_budget_reset_at": null, + "user_api_key_team_id": null, + "user_api_key_org_id": null, + "user_api_key_org_alias": null, + "user_api_key_project_id": null, + "user_api_key_project_alias": null, + "user_api_key_user_id": "default_user_id", + "user_api_key_team_alias": null, + "user_api_key_user_email": null, + "user_api_key_end_user_id": null, + "user_api_key_request_route": "/v1/chat/completions", + "spend_logs_metadata": null, + "requester_ip_address": "127.0.0.1", + "user_agent": "PostmanRuntime/7.53.0", + "requester_metadata": { + "headers": { + "content-type": "application/json", + "user-agent": "PostmanRuntime/7.53.0", + "accept": "*/*", + "postman-token": "08d81a08-addf-4a27-a9f8-302f317d7296", + "host": "0.0.0.0:4000", + "accept-encoding": "gzip, deflate, br", + "connection": "keep-alive", + "content-length": "183" + } + }, + "prompt_management_metadata": null, + "applied_guardrails": [], + "mcp_tool_call_metadata": null, + "vector_store_request_metadata": null, + "usage_object": { + "completion_tokens": 76, + "prompt_tokens": 19, + "total_tokens": 95, + "completion_tokens_details": { + "reasoning_tokens": 0, + "text_tokens": 76 + }, + "prompt_tokens_details": { + "audio_tokens": null, + "cached_tokens": 0, + "text_tokens": 19, + "image_tokens": null, + "video_tokens": null, + "cache_creation_tokens": 0, + "cache_creation_token_details": { + "ephemeral_5m_input_tokens": 0, + "ephemeral_1h_input_tokens": 0 + } + }, + "cache_creation_input_tokens": 0, + "cache_read_input_tokens": 0, + "inference_geo": "global", + "speed": null + }, + "requester_custom_headers": {}, + "cold_storage_object_key": null, + "user_api_key_auth_metadata": {}, + "team_alias": null, + "team_id": null + }, + "cache_key": null, + "response_cost": 0.001995, + "cost_breakdown": { + "input_cost": 9.5e-05, + "output_cost": 0.0019, + "total_cost": 0.001995, + "tool_usage_cost": 0.0, + "original_cost": 0.001995, + "discount_percent": 0.0, + "discount_amount": 0.0, + "margin_percent": 0.0, + "margin_fixed_amount": 0.0, + "margin_total_amount": 0.0 + }, + "total_tokens": 95, + "prompt_tokens": 19, + "completion_tokens": 76, + "request_tags": [ + "User-Agent: PostmanRuntime", + "User-Agent: PostmanRuntime/7.53.0" + ], + "end_user": "", + "api_base": "https://api.anthropic.com/v1/messages", + "model_group": "claude-opus-4-7", + "model_id": "91a8591ffdee6edec7016031dc8a092b14ce59f85b1b105c0afa69280775d380", + "requester_ip_address": "127.0.0.1", + "user_agent": "PostmanRuntime/7.53.0", + "messages": [ + { + "role": "user", + "content": "Hello, how are you?" + } + ], + "response": { + "id": "chatcmpl-db0b1b06-c051-4281-98ff-6add4a98effb", + "created": 1776686262, + "model": "claude-opus-4-7", + "object": "chat.completion", + "system_fingerprint": null, + "choices": [ + { + "finish_reason": "stop", + "index": 0, + "message": { + "content": "Hello! I'm doing well, thanks for asking. I'm ready to help with whatever you need\u2014whether that's answering questions, brainstorming ideas, working through a problem, or just having a conversation. What can I do for you today?", + "role": "assistant", + "tool_calls": null, + "function_call": null, + "provider_specific_fields": { + "citations": null, + "thinking_blocks": null + } + } + } + ], + "usage": { + "completion_tokens": 76, + "prompt_tokens": 19, + "total_tokens": 95, + "completion_tokens_details": { + "reasoning_tokens": 0, + "text_tokens": 76 + }, + "prompt_tokens_details": { + "audio_tokens": null, + "cached_tokens": 0, + "text_tokens": 19, + "image_tokens": null, + "video_tokens": null, + "cache_creation_tokens": 0, + "cache_creation_token_details": { + "ephemeral_5m_input_tokens": 0, + "ephemeral_1h_input_tokens": 0 + } + }, + "cache_creation_input_tokens": 0, + "cache_read_input_tokens": 0, + "inference_geo": "global", + "speed": null + } + }, + "model_parameters": { + "stream": false, + "reasoning_effort": "max" + }, + "hidden_params": { + "model_id": "91a8591ffdee6edec7016031dc8a092b14ce59f85b1b105c0afa69280775d380", + "cache_key": null, + "api_base": null, + "response_cost": 0.001995, + "additional_headers": { + "x_ratelimit_limit_requests": 50, + "x_ratelimit_limit_tokens": 38000, + "x_ratelimit_remaining_requests": 49, + "x_ratelimit_remaining_tokens": 38000, + "llm_provider-date": "Mon, 20 Apr 2026 11:57:42 GMT", + "llm_provider-content-type": "application/json", + "llm_provider-transfer-encoding": "chunked", + "llm_provider-connection": "keep-alive", + "llm_provider-anthropic-ratelimit-input-tokens-limit": "30000", + "llm_provider-anthropic-ratelimit-input-tokens-remaining": "30000", + "llm_provider-anthropic-ratelimit-input-tokens-reset": "2026-04-20T11:57:41Z", + "llm_provider-anthropic-ratelimit-output-tokens-limit": "8000", + "llm_provider-anthropic-ratelimit-output-tokens-remaining": "8000", + "llm_provider-anthropic-ratelimit-output-tokens-reset": "2026-04-20T11:57:42Z", + "llm_provider-anthropic-ratelimit-requests-limit": "50", + "llm_provider-anthropic-ratelimit-requests-remaining": "49", + "llm_provider-anthropic-ratelimit-requests-reset": "2026-04-20T11:57:41Z", + "llm_provider-anthropic-ratelimit-tokens-limit": "38000", + "llm_provider-anthropic-ratelimit-tokens-remaining": "38000", + "llm_provider-anthropic-ratelimit-tokens-reset": "2026-04-20T11:57:41Z", + "llm_provider-request-id": "req_011CaEyx63z2yDaDF1967jyU", + "llm_provider-strict-transport-security": "max-age=31536000; includeSubDomains; preload", + "llm_provider-anthropic-organization-id": "cf4012a8-5914-4165-9573-d00b88e3d39e", + "llm_provider-server": "cloudflare", + "llm_provider-x-envoy-upstream-service-time": "2017", + "llm_provider-content-encoding": "gzip", + "llm_provider-vary": "Accept-Encoding", + "llm_provider-server-timing": "x-originResponse;dur=2019", + "llm_provider-cf-cache-status": "DYNAMIC", + "llm_provider-set-cookie": "_cfuvid=tNQ_aMy8wM6wZvYT8Gx.fxKqkk22mUq7vp0IzBG.r8g-1776686259.790975-1.0.1.1-eRorMbMzA43iO9c9lvJ8g5QwV1nQi.QPuE8FLWcVVeM; HttpOnly; SameSite=None; Secure; Path=/; Domain=api.anthropic.com", + "llm_provider-content-security-policy": "default-src 'none'; frame-ancestors 'none'", + "llm_provider-x-robots-tag": "none", + "llm_provider-cf-ray": "9ef3f903ab1347e7-BOM", + "llm_provider-x-ratelimit-limit-requests": "50", + "llm_provider-x-ratelimit-remaining-requests": "49", + "llm_provider-x-ratelimit-limit-tokens": "38000", + "llm_provider-x-ratelimit-remaining-tokens": "38000", + "x-litellm-model-group": "claude-opus-4-7", + "x-litellm-attempted-retries": 0, + "x-litellm-attempted-fallbacks": 0 + }, + "litellm_overhead_time_ms": 15.177, + "batch_models": null, + "litellm_model_name": "anthropic/claude-opus-4-7", + "usage_object": null + }, + "model_map_information": { + "model_map_key": "claude-opus-4-7", + "model_map_value": { + "key": "claude-opus-4-7", + "max_tokens": 128000, + "max_input_tokens": 1000000, + "max_output_tokens": 128000, + "input_cost_per_token": 5e-06, + "input_cost_per_token_flex": null, + "input_cost_per_token_priority": null, + "cache_creation_input_token_cost": 6.25e-06, + "cache_creation_input_token_cost_above_200k_tokens": null, + "cache_read_input_token_cost": 5e-07, + "cache_read_input_token_cost_above_200k_tokens": null, + "cache_read_input_token_cost_above_272k_tokens": null, + "cache_read_input_token_cost_flex": null, + "cache_read_input_token_cost_priority": null, + "cache_creation_input_token_cost_above_1hr": 1e-05, + "input_cost_per_character": null, + "input_cost_per_token_above_128k_tokens": null, + "input_cost_per_token_above_200k_tokens": null, + "input_cost_per_token_above_272k_tokens": null, + "input_cost_per_query": null, + "input_cost_per_second": null, + "input_cost_per_audio_token": null, + "input_cost_per_image_token": null, + "input_cost_per_image": null, + "input_cost_per_audio_per_second": null, + "input_cost_per_video_per_second": null, + "input_cost_per_token_batches": null, + "output_cost_per_token_batches": null, + "output_cost_per_token": 2.5e-05, + "output_cost_per_token_flex": null, + "output_cost_per_token_priority": null, + "output_cost_per_audio_token": null, + "output_cost_per_character": null, + "output_cost_per_reasoning_token": null, + "output_cost_per_token_above_128k_tokens": null, + "output_cost_per_character_above_128k_tokens": null, + "output_cost_per_token_above_200k_tokens": null, + "output_cost_per_token_above_272k_tokens": null, + "output_cost_per_second": null, + "output_cost_per_second_1080p": null, + "output_cost_per_video_per_second": null, + "output_cost_per_image": null, + "output_cost_per_image_token": null, + "output_vector_size": null, + "citation_cost_per_token": null, + "tiered_pricing": null, + "litellm_provider": "anthropic", + "mode": "chat", + "supports_system_messages": null, + "supports_response_schema": true, + "supports_vision": true, + "supports_function_calling": true, + "supports_tool_choice": true, + "supports_assistant_prefill": false, + "supports_prompt_caching": true, + "supports_audio_input": null, + "supports_audio_output": null, + "supports_pdf_input": true, + "supports_embedding_image_input": null, + "supports_native_streaming": null, + "supports_native_structured_output": null, + "supports_web_search": null, + "supports_url_context": null, + "supports_reasoning": true, + "supports_none_reasoning_effort": null, + "supports_xhigh_reasoning_effort": true, + "supports_computer_use": true, + "search_context_cost_per_query": { + "search_context_size_high": 0.01, + "search_context_size_low": 0.01, + "search_context_size_medium": 0.01 + }, + "tpm": null, + "rpm": null, + "ocr_cost_per_page": null, + "annotation_cost_per_page": null, + "provider_specific_entry": { + "us": 1.1, + "fast": 6.0 + }, + "uses_embed_content": null, + "supported_openai_params": [ + "stream", + "stop", + "temperature", + "top_p", + "max_tokens", + "max_completion_tokens", + "tools", + "tool_choice", + "extra_headers", + "parallel_tool_calls", + "response_format", + "user", + "web_search_options", + "speed", + "context_management", + "cache_control", + "thinking", + "reasoning_effort" + ] + } + }, + "error_str": null, + "error_information": { + "error_code": "", + "error_class": "", + "llm_provider": "", + "traceback": "", + "error_message": "" + }, + "response_cost_failure_debug_info": null, + "guardrail_information": null, + "standard_built_in_tools_params": { + "web_search_options": null, + "file_search": null + } +}17:27:42 - LiteLLM:DEBUG: callback_controls.py:37 - Dynamically disabled callbacks from x-litellm-disable-callbacks: None +17:27:42 - LiteLLM:DEBUG: callback_controls.py:38 - Checking if > is disabled via headers. Disable callbacks from headers: None +17:27:42 - LiteLLM:DEBUG: utils.py:5620 - checking potential_model_names in litellm.model_cost: {'split_model': 'claude-opus-4-7', 'combined_model_name': 'anthropic/claude-opus-4-7', 'stripped_model_name': 'anthropic/claude-opus-4-7', 'combined_stripped_model_name': 'anthropic/claude-opus-4-7', 'custom_llm_provider': 'anthropic'} +17:27:42 - LiteLLM:DEBUG: callback_controls.py:37 - Dynamically disabled callbacks from x-litellm-disable-callbacks: None +17:27:42 - LiteLLM:DEBUG: callback_controls.py:38 - Checking if is disabled via headers. Disable callbacks from headers: None +17:27:42 - LiteLLM Proxy:DEBUG: proxy_track_cost_callback.py:152 - INSIDE _PROXY_track_cost_callback +17:27:42 - LiteLLM Proxy:DEBUG: proxy_track_cost_callback.py:154 - kwargs stream: False + complete streaming response: None +17:27:42 - LiteLLM Proxy:DEBUG: proxy_track_cost_callback.py:186 - user_api_key 88dc28d0f030c55ed4ab77ed8faf098196cb1c05df778539800c9f1243fe6b4b, user_id default_user_id, team_id None, end_user_id None +17:27:42 - LiteLLM Proxy:DEBUG: db_spend_update_writer.py:119 - Enters prisma db call, response_cost: 0.001995, token: 88dc28d0f030c55ed4ab77ed8faf098196cb1c05df778539800c9f1243fe6b4b; user_id: default_user_id; team_id: None +17:27:42 - LiteLLM Proxy:DEBUG: spend_tracking_utils.py:116 - getting payload for SpendLogs, available keys in metadata: ['headers', 'requester_metadata', 'user_api_key_hash', 'user_api_key_alias', 'user_api_key_spend', 'user_api_key_max_budget', 'user_api_key_team_id', 'user_api_key_project_id', 'user_api_key_project_alias', 'user_api_key_user_id', 'user_api_key_org_id', 'user_api_key_org_alias', 'user_api_key_team_alias', 'user_api_key_end_user_id', 'user_api_key_user_email', 'user_api_key_request_route', 'user_api_key_budget_reset_at', 'user_api_key_auth_metadata', 'user_api_key', 'agent_id', 'user_api_end_user_max_budget', 'user_api_key_auth', 'litellm_api_version', 'global_max_parallel_requests', 'user_api_key_team_max_budget', 'user_api_key_team_spend', 'user_api_key_model_max_budget', 'user_api_key_end_user_model_max_budget', 'user_api_key_user_spend', 'user_api_key_user_max_budget', 'user_api_key_metadata', 'user_api_key_team_metadata', 'user_api_key_object_permission_id', 'user_api_key_team_object_permission_id', 'endpoint', 'litellm_parent_otel_span', 'requester_ip_address', 'user_agent', 'queue_time_seconds', 'model_group', 'model_group_alias', 'model_group_size', 'attempted_retries', 'max_retries', 'deployment', 'model_info', 'api_base', 'deployment_model_name', 'caching_groups', 'hidden_params'] +17:27:42 - LiteLLM Proxy:DEBUG: spend_tracking_utils.py:480 - SpendTable: created payload - request_id: chatcmpl-db0b1b06-c051-4281-98ff-6add4a98effb, model: anthropic/claude-opus-4-7, spend: 0.001995 +17:27:42 - LiteLLM Proxy:DEBUG: db_spend_update_writer.py:726 - Writing spend log to db - request_id: chatcmpl-db0b1b06-c051-4281-98ff-6add4a98effb, spend: 0.001995 +17:27:42 - LiteLLM Proxy:DEBUG: db_spend_update_writer.py:193 - Runs spend update on all tables +17:27:42 - LiteLLM:DEBUG: callback_controls.py:37 - Dynamically disabled callbacks from x-litellm-disable-callbacks: None +17:27:42 - LiteLLM:DEBUG: callback_controls.py:38 - Checking if is disabled via headers. Disable callbacks from headers: None +17:27:42 - LiteLLM:DEBUG: callback_controls.py:37 - Dynamically disabled callbacks from x-litellm-disable-callbacks: None +17:27:42 - LiteLLM:DEBUG: callback_controls.py:38 - Checking if is disabled via headers. Disable callbacks from headers: None +17:27:42 - LiteLLM Proxy:DEBUG: model_max_budget_limiter.py:227 - in RouterBudgetLimiting.async_log_success_event +17:27:42 - LiteLLM Proxy:DEBUG: model_max_budget_limiter.py:252 - Not running _PROXY_VirtualKeyModelMaxBudgetLimiter.async_log_success_event because user_api_key_model_max_budget and user_api_key_end_user_model_max_budget are None or empty. +17:27:42 - LiteLLM:DEBUG: callback_controls.py:37 - Dynamically disabled callbacks from x-litellm-disable-callbacks: None +17:27:42 - LiteLLM:DEBUG: callback_controls.py:38 - Checking if is disabled via headers. Disable callbacks from headers: None +17:27:42 - LiteLLM:DEBUG: callback_controls.py:37 - Dynamically disabled callbacks from x-litellm-disable-callbacks: None +17:27:42 - LiteLLM:DEBUG: callback_controls.py:38 - Checking if is disabled via headers. Disable callbacks from headers: None +17:27:42 - LiteLLM Proxy:DEBUG: parallel_request_limiter_v3.py:1740 - INSIDE parallel request limiter ASYNC SUCCESS LOGGING +17:27:42 - LiteLLM Proxy:DEBUG: parallel_request_limiter_v3.py:1506 - TTL preservation script not available, using regular pipeline +17:27:42 - LiteLLM:DEBUG: callback_controls.py:37 - Dynamically disabled callbacks from x-litellm-disable-callbacks: None +17:27:42 - LiteLLM:DEBUG: callback_controls.py:38 - Checking if is disabled via headers. Disable callbacks from headers: None +17:27:42 - LiteLLM:DEBUG: callback_controls.py:37 - Dynamically disabled callbacks from x-litellm-disable-callbacks: None +17:27:42 - LiteLLM:DEBUG: callback_controls.py:38 - Checking if is disabled via headers. Disable callbacks from headers: None +17:27:42 - LiteLLM:DEBUG: callback_controls.py:37 - Dynamically disabled callbacks from x-litellm-disable-callbacks: None +17:27:42 - LiteLLM:DEBUG: callback_controls.py:38 - Checking if is disabled via headers. Disable callbacks from headers: None +17:27:42 - LiteLLM:DEBUG: callback_controls.py:37 - Dynamically disabled callbacks from x-litellm-disable-callbacks: None +17:27:42 - LiteLLM:DEBUG: callback_controls.py:38 - Checking if is disabled via headers. Disable callbacks from headers: None +17:27:42 - LiteLLM:DEBUG: callback_controls.py:37 - Dynamically disabled callbacks from x-litellm-disable-callbacks: None +17:27:42 - LiteLLM:DEBUG: callback_controls.py:38 - Checking if is disabled via headers. Disable callbacks from headers: None +17:27:42 - LiteLLM:DEBUG: callback_controls.py:37 - Dynamically disabled callbacks from x-litellm-disable-callbacks: None +17:27:42 - LiteLLM:DEBUG: callback_controls.py:38 - Checking if is disabled via headers. Disable callbacks from headers: None +17:27:42 - LiteLLM:DEBUG: callback_controls.py:37 - Dynamically disabled callbacks from x-litellm-disable-callbacks: None +17:27:42 - LiteLLM:DEBUG: callback_controls.py:38 - Checking if is disabled via headers. Disable callbacks from headers: None +17:27:42 - LiteLLM Proxy:DEBUG: spend_update_queue.py:44 - Adding update to queue: {'entity_type': , 'entity_id': 'default_user_id', 'response_cost': 0.001995} +17:27:42 - LiteLLM Proxy:DEBUG: spend_update_queue.py:44 - Adding update to queue: {'entity_type': , 'entity_id': '88dc28d0f030c55ed4ab77ed8faf098196cb1c05df778539800c9f1243fe6b4b', 'response_cost': 0.001995} +17:27:42 - LiteLLM Proxy:DEBUG: db_spend_update_writer.py:562 - track_cost_callback: team_id is None or prisma_client is None. Not tracking spend for team +17:27:42 - LiteLLM Proxy:DEBUG: db_spend_update_writer.py:616 - track_cost_callback: org_id is None or prisma_client is None. Not tracking spend for org +17:27:42 - LiteLLM Proxy:DEBUG: spend_update_queue.py:44 - Adding update to queue: {'entity_type': , 'entity_id': 'User-Agent: PostmanRuntime', 'response_cost': 0.001995} +17:27:42 - LiteLLM Proxy:DEBUG: spend_update_queue.py:44 - Adding update to queue: {'entity_type': , 'entity_id': 'User-Agent: PostmanRuntime/7.53.0', 'response_cost': 0.001995} +17:27:42 - LiteLLM Proxy:DEBUG: db_spend_update_writer.py:1950 - Logged request status: success +17:27:42 - LiteLLM Proxy:DEBUG: daily_spend_update_queue.py:63 - Adding update to queue: {'default_user_id_2026-04-20_88dc28d0f030c55ed4ab77ed8faf098196cb1c05df778539800c9f1243fe6b4b_anthropic/claude-opus-4-7_anthropic_/chat/completions': {'user_id': 'default_user_id', 'date': '2026-04-20', 'REDACTED', 'model': 'anthropic/claude-opus-4-7', 'model_group': 'claude-opus-4-7', 'mcp_namespaced_tool_name': None, 'custom_llm_provider': 'anthropic', 'endpoint': '/chat/completions', 'prompt_tokens': 19, 'completion_tokens': 76, 'spend': 0.001995, 'api_requests': 1, 'successful_requests': 1, 'failed_requests': 0, 'cache_read_input_tokens': 0, 'cache_creation_input_tokens': 0}} +17:27:42 - LiteLLM Proxy:DEBUG: db_spend_update_writer.py:2119 - end_user is None or empty for request. Skipping incrementing end user spend. +17:27:42 - LiteLLM Proxy:DEBUG: db_spend_update_writer.py:1950 - Logged request status: success +17:27:42 - LiteLLM Proxy:DEBUG: daily_spend_update_queue.py:63 - Adding update to queue: {'_2026-04-20_88dc28d0f030c55ed4ab77ed8faf098196cb1c05df778539800c9f1243fe6b4b_anthropic/claude-opus-4-7_anthropic_/chat/completions': {'team_id': '', 'date': '2026-04-20', 'REDACTED', 'model': 'anthropic/claude-opus-4-7', 'model_group': 'claude-opus-4-7', 'mcp_namespaced_tool_name': None, 'custom_llm_provider': 'anthropic', 'endpoint': '/chat/completions', 'prompt_tokens': 19, 'completion_tokens': 76, 'spend': 0.001995, 'api_requests': 1, 'successful_requests': 1, 'failed_requests': 0, 'cache_read_input_tokens': 0, 'cache_creation_input_tokens': 0}} +17:27:42 - LiteLLM Proxy:DEBUG: db_spend_update_writer.py:2076 - organization_id is None for request. Skipping incrementing organization spend. +17:27:42 - LiteLLM Proxy:DEBUG: db_spend_update_writer.py:1950 - Logged request status: success +17:27:42 - LiteLLM Proxy:DEBUG: daily_spend_update_queue.py:63 - Adding update to queue: {'User-Agent: PostmanRuntime_2026-04-20_88dc28d0f030c55ed4ab77ed8faf098196cb1c05df778539800c9f1243fe6b4b_anthropic/claude-opus-4-7_anthropic_/chat/completions': {'tag': 'User-Agent: PostmanRuntime', 'date': '2026-04-20', 'REDACTED', 'model': 'anthropic/claude-opus-4-7', 'model_group': 'claude-opus-4-7', 'mcp_namespaced_tool_name': None, 'custom_llm_provider': 'anthropic', 'endpoint': '/chat/completions', 'prompt_tokens': 19, 'completion_tokens': 76, 'spend': 0.001995, 'api_requests': 1, 'successful_requests': 1, 'failed_requests': 0, 'cache_read_input_tokens': 0, 'cache_creation_input_tokens': 0, 'request_id': 'chatcmpl-db0b1b06-c051-4281-98ff-6add4a98effb'}} +17:27:42 - LiteLLM Proxy:DEBUG: daily_spend_update_queue.py:63 - Adding update to queue: {'User-Agent: PostmanRuntime/7.53.0_2026-04-20_88dc28d0f030c55ed4ab77ed8faf098196cb1c05df778539800c9f1243fe6b4b_anthropic/claude-opus-4-7_anthropic_/chat/completions': {'tag': 'User-Agent: PostmanRuntime/7.53.0', 'date': '2026-04-20', 'REDACTED', 'model': 'anthropic/claude-opus-4-7', 'model_group': 'claude-opus-4-7', 'mcp_namespaced_tool_name': None, 'custom_llm_provider': 'anthropic', 'endpoint': '/chat/completions', 'prompt_tokens': 19, 'completion_tokens': 76, 'spend': 0.001995, 'api_requests': 1, 'successful_requests': 1, 'failed_requests': 0, 'cache_read_input_tokens': 0, 'cache_creation_input_tokens': 0, 'request_id': 'chatcmpl-db0b1b06-c051-4281-98ff-6add4a98effb'}} +17:27:42 - LiteLLM Proxy:DEBUG: proxy_server.py:1945 - _update_key_cache: hashed_token=88dc28d0f030c55ed4ab77ed8faf098196cb1c05df778539800c9f1243fe6b4b +17:27:42 - LiteLLM Proxy:DEBUG: proxy_server.py:1947 - _update_key_cache: existing_spend_obj=token='88dc28d0f030c55ed4ab77ed8faf098196cb1c05df778539800c9f1243fe6b4b' key_name=None key_alias=None spend=0.0 max_budget=None expires=None models=[] aliases={} config={} user_id='default_user_id' team_id=None agent_id=None project_id=None max_parallel_requests=None metadata={} tpm_limit=None rpm_limit=None budget_duration=None budget_reset_at=None allowed_cache_controls=[] allowed_routes=[] permissions={} model_spend={} model_max_budget={} soft_budget_cooldown=False blocked=None litellm_budget_table=None org_id=None created_at=None created_by=None updated_at=None updated_by=None last_active=None object_permission_id=None object_permission=None access_group_ids=None rotation_count=0 auto_rotate=False rotation_interval=None last_rotation_at=None key_rotation_at=None router_settings=None budget_limits=None team_spend=None team_alias=None team_tpm_limit=None team_rpm_limit=None team_max_budget=None team_soft_budget=None team_models=[] team_blocked=False soft_budget=None team_model_aliases=None team_member=None team_metadata=None team_object_permission_id=None team_member_spend=None team_member_tpm_limit=None team_member_rpm_limit=None end_user_id=None end_user_tpm_limit=None end_user_rpm_limit=None end_user_max_budget=None end_user_model_max_budget=None organization_alias=None organization_max_budget=None organization_tpm_limit=None organization_rpm_limit=None organization_metadata=None project_alias=None project_metadata=None last_refreshed_at=1776686259.698418 REDACTED' user_role= allowed_model_region=None parent_otel_span=None rpm_limit_per_model=None tpm_limit_per_model=None user_tpm_limit=None user_rpm_limit=None user_email=None user_spend=None user_max_budget=None request_route='/v1/chat/completions' user=None created_by_user=None end_user_object_permission=None jwt_claims=None +17:27:47 - LiteLLM Proxy:DEBUG: utils.py:5096 - Spend logs queue size (1) below threshold (100), processing with backoff +17:27:47 - LiteLLM Proxy:INFO: utils.py:4813 - Spend tracking - processing 1 spend logs for DB write +17:27:47 - LiteLLM Proxy:DEBUG: utils.py:4850 - Flushed 1 logs to the DB. +17:27:47 - LiteLLM Proxy:DEBUG: utils.py:4859 - 1 logs processed. Remaining in queue: 0 +17:27:47 - LiteLLM Proxy:INFO: spend_update_queue.py:35 - Spend tracking - flushed 4 spend update items from in-memory queue +17:27:47 - LiteLLM Proxy:DEBUG: spend_update_queue.py:39 - Aggregating updates by entity type: [{'entity_type': , 'entity_id': 'default_user_id', 'response_cost': 0.001995}, {'entity_type': , 'entity_id': '88dc28d0f030c55ed4ab77ed8faf098196cb1c05df778539800c9f1243fe6b4b', 'response_cost': 0.001995}, {'entity_type': , 'entity_id': 'User-Agent: PostmanRuntime', 'response_cost': 0.001995}, {'entity_type': , 'entity_id': 'User-Agent: PostmanRuntime/7.53.0', 'response_cost': 0.001995}] +17:27:47 - LiteLLM Proxy:DEBUG: db_spend_update_writer.py:1122 - User Spend transactions: {'default_user_id': 0.001995} +17:27:47 - LiteLLM Proxy:DEBUG: db_spend_update_writer.py:1165 - End-User Spend transactions: {} +17:27:47 - LiteLLM Proxy:DEBUG: db_spend_update_writer.py:1180 - KEY Spend transactions: {'88dc28d0f030c55ed4ab77ed8faf098196cb1c05df778539800c9f1243fe6b4b': 0.001995} +17:27:47 - LiteLLM Proxy:DEBUG: db_spend_update_writer.py:1221 - Team Spend transactions: {} +17:27:47 - LiteLLM Proxy:DEBUG: db_spend_update_writer.py:1269 - Team Membership Spend transactions: {} +17:27:47 - LiteLLM Proxy:DEBUG: db_spend_update_writer.py:1339 - Org Spend transactions: {} +17:27:47 - LiteLLM Proxy:DEBUG: db_spend_update_writer.py:1432 - Tag Spend transactions: {'User-Agent: PostmanRuntime': 0.001995, 'User-Agent: PostmanRuntime/7.53.0': 0.001995} +17:27:47 - LiteLLM Proxy:DEBUG: db_spend_update_writer.py:1442 - Updating spend for Tag tag_name=User-Agent: PostmanRuntime by 0.001995 +17:27:47 - LiteLLM Proxy:DEBUG: db_spend_update_writer.py:1442 - Updating spend for Tag tag_name=User-Agent: PostmanRuntime/7.53.0 by 0.001995 +17:27:47 - LiteLLM Proxy:DEBUG: db_spend_update_writer.py:1432 - Agent Spend transactions: {} +17:27:47 - LiteLLM Proxy:INFO: daily_spend_update_queue.py:90 - Spend tracking - flushed 1 daily spend update items from in-memory queue +17:27:47 - LiteLLM Proxy:DEBUG: daily_spend_update_queue.py:99 - Aggregated daily spend update transactions: {'default_user_id_2026-04-20_88dc28d0f030c55ed4ab77ed8faf098196cb1c05df778539800c9f1243fe6b4b_anthropic/claude-opus-4-7_anthropic_/chat/completions': {'user_id': 'default_user_id', 'date': '2026-04-20', 'REDACTED', 'model': 'anthropic/claude-opus-4-7', 'model_group': 'claude-opus-4-7', 'mcp_namespaced_tool_name': None, 'custom_llm_provider': 'anthropic', 'endpoint': '/chat/completions', 'prompt_tokens': 19, 'completion_tokens': 76, 'spend': 0.001995, 'api_requests': 1, 'successful_requests': 1, 'failed_requests': 0, 'cache_read_input_tokens': 0, 'cache_creation_input_tokens': 0}} +17:27:47 - LiteLLM Proxy:DEBUG: db_spend_update_writer.py:1573 - Daily User Spend transactions: 1 +17:27:47 - LiteLLM Proxy:DEBUG: db_spend_update_writer.py:1744 - Processed 1 daily user transactions in 0.02s +17:27:47 - LiteLLM Proxy:INFO: daily_spend_update_queue.py:90 - Spend tracking - flushed 1 daily spend update items from in-memory queue +17:27:47 - LiteLLM Proxy:DEBUG: daily_spend_update_queue.py:99 - Aggregated daily spend update transactions: {'_2026-04-20_88dc28d0f030c55ed4ab77ed8faf098196cb1c05df778539800c9f1243fe6b4b_anthropic/claude-opus-4-7_anthropic_/chat/completions': {'team_id': '', 'date': '2026-04-20', 'REDACTED', 'model': 'anthropic/claude-opus-4-7', 'model_group': 'claude-opus-4-7', 'mcp_namespaced_tool_name': None, 'custom_llm_provider': 'anthropic', 'endpoint': '/chat/completions', 'prompt_tokens': 19, 'completion_tokens': 76, 'spend': 0.001995, 'api_requests': 1, 'successful_requests': 1, 'failed_requests': 0, 'cache_read_input_tokens': 0, 'cache_creation_input_tokens': 0}} +17:27:47 - LiteLLM Proxy:DEBUG: db_spend_update_writer.py:1573 - Daily Team Spend transactions: 1 +17:27:47 - LiteLLM Proxy:DEBUG: db_spend_update_writer.py:1744 - Processed 1 daily team transactions in 0.01s +17:27:47 - LiteLLM Proxy:DEBUG: daily_spend_update_queue.py:99 - Aggregated daily spend update transactions: {} +17:27:47 - LiteLLM Proxy:DEBUG: db_spend_update_writer.py:1573 - Daily Org Spend transactions: 0 +17:27:47 - LiteLLM Proxy:DEBUG: db_spend_update_writer.py:1606 - No new transactions to process for daily org spend update +17:27:47 - LiteLLM Proxy:DEBUG: daily_spend_update_queue.py:99 - Aggregated daily spend update transactions: {} +17:27:47 - LiteLLM Proxy:DEBUG: db_spend_update_writer.py:1573 - Daily End_user Spend transactions: 0 +17:27:47 - LiteLLM Proxy:DEBUG: db_spend_update_writer.py:1606 - No new transactions to process for daily end_user spend update +17:27:47 - LiteLLM Proxy:DEBUG: daily_spend_update_queue.py:99 - Aggregated daily spend update transactions: {} +17:27:47 - LiteLLM Proxy:DEBUG: db_spend_update_writer.py:1573 - Daily Agent Spend transactions: 0 +17:27:47 - LiteLLM Proxy:DEBUG: db_spend_update_writer.py:1606 - No new transactions to process for daily agent spend update +17:27:47 - LiteLLM Proxy:DEBUG: utils.py:4931 - Spend Logs transactions: 0 +17:28:01 - LiteLLM Proxy:ERROR: utils.py:3918 - prisma-query-engine PID 66228 exited (waitpid thread); triggering reconnect. +17:28:01 - LiteLLM Proxy:WARNING: utils.py:4178 - Attempting Prisma DB reconnect. reason=engine_process_death +17:28:01 - LiteLLM Proxy:WARNING: utils.py:4107 - prisma-query-engine PID 0 is dead; reconnecting. +query-engine ac9d7041ed77bcc8a8dbd2ab6616b39013829574 +INFO: Shutting down +INFO: Waiting for application shutdown. +17:28:01 - LiteLLM Proxy:INFO: proxy_server.py:971 - SESSION REUSE: Closed shared aiohttp session +17:28:01 - LiteLLM Proxy:DEBUG: utils.py:4082 - Stopped engine process watcher. +17:28:01 - LiteLLM Proxy:INFO: utils.py:4322 - Stopped Prisma DB health watchdog +17:28:01 - LiteLLM Proxy:INFO: proxy_server.py:708 - Shutting down LiteLLM Proxy Server +17:28:01 - LiteLLM Proxy:DEBUG: proxy_server.py:710 - Disconnecting from Prisma +INFO: Application shutdown complete. +INFO: Finished server process [66039] +17:28:01 - LiteLLM:DEBUG: logging_worker.py:134 - LoggingWorker cancelled during shutdown +17:28:01 - LiteLLM:DEBUG: logging_worker.py:446 - [LoggingWorker] atexit: Queue is empty +