diff --git a/litellm/model_prices_and_context_window_backup.json b/litellm/model_prices_and_context_window_backup.json index 14d3c1ae846..195e2885770 100644 --- a/litellm/model_prices_and_context_window_backup.json +++ b/litellm/model_prices_and_context_window_backup.json @@ -9084,9 +9084,6 @@ "cache_read_input_token_cost": 5e-07, "input_cost_per_token": 5e-06, "litellm_provider": "anthropic", - "aliases": [ - "claude-opus-4-6-20250514" - ], "max_input_tokens": 1000000, "max_output_tokens": 128000, "max_tokens": 128000, diff --git a/proxy_server.log b/proxy_server.log deleted file mode 100644 index 6aae2d1bcbf..00000000000 --- a/proxy_server.log +++ /dev/null @@ -1,659 +0,0 @@ -:128: RuntimeWarning: 'litellm.proxy.proxy_cli' found in sys.modules after import of package 'litellm.proxy', but prior to execution of 'litellm.proxy.proxy_cli'; this may result in unpredictable behaviour -17:27:32 - LiteLLM:ERROR: check_migration.py:100 - 🚨🚨🚨 prisma schema out of sync with db. Consider running these sql_commands to sync the two - ['ALTER TABLE "LiteLLM_BudgetTable" ADD COLUMN "allowed_models" TEXT[] DEFAULT ARRAY[]::TEXT[];', 'ALTER TABLE "LiteLLM_MCPServerTable" ADD COLUMN "instructions" TEXT;', 'ALTER TABLE "LiteLLM_TeamTable" ADD COLUMN "budget_limits" JSONB, ADD COLUMN "default_team_member_models" TEXT[] DEFAULT ARRAY[]::TEXT[];', 'ALTER TABLE "LiteLLM_VerificationToken" ADD COLUMN "budget_limits" JSONB;', 'CREATE INDEX "LiteLLM_HealthCheckTable_model_id_model_name_checked_at_idx" ON "LiteLLM_HealthCheckTable"("model_id", "model_name", "checked_at" DESC);'] -NoneType: None -INFO: Started server process [66039] -INFO: Waiting for application startup. -17:27:32 - LiteLLM Proxy:DEBUG: proxy_server.py:804 - litellm.proxy.proxy_server.py::startup() - CHECKING PREMIUM USER - True -17:27:32 - LiteLLM Proxy:DEBUG: proxy_server.py:817 - worker_config: {"model": null, "alias": null, "api_base": null, "api_version": "2025-02-01-preview", "debug": false, "detailed_debug": true, "temperature": null, "max_tokens": null, "request_timeout": null, "max_budget": null, "telemetry": true, "drop_params": false, "add_function_to_prompt": false, "headers": null, "save": false, "config": "proxy_server_config.yaml", "use_queue": false} -Changes to DB Schema detected -Required SQL commands: -ALTER TABLE "LiteLLM_BudgetTable" ADD COLUMN "allowed_models" TEXT[] DEFAULT ARRAY[]::TEXT[]; -ALTER TABLE "LiteLLM_MCPServerTable" ADD COLUMN "instructions" TEXT; -ALTER TABLE "LiteLLM_TeamTable" ADD COLUMN "budget_limits" JSONB, ADD COLUMN "default_team_member_models" TEXT[] DEFAULT ARRAY[]::TEXT[]; -ALTER TABLE "LiteLLM_VerificationToken" ADD COLUMN "budget_limits" JSONB; -CREATE INDEX "LiteLLM_HealthCheckTable_model_id_model_name_checked_at_idx" ON "LiteLLM_HealthCheckTable"("model_id", "model_name", "checked_at" DESC); - - β–ˆβ–ˆβ•— β–ˆβ–ˆβ•—β–ˆβ–ˆβ–ˆβ–ˆβ–ˆβ–ˆβ–ˆβ–ˆβ•—β–ˆβ–ˆβ–ˆβ–ˆβ–ˆβ–ˆβ–ˆβ•—β–ˆβ–ˆβ•— β–ˆβ–ˆβ•— β–ˆβ–ˆβ–ˆβ•— β–ˆβ–ˆβ–ˆβ•— - β–ˆβ–ˆβ•‘ β–ˆβ–ˆβ•‘β•šβ•β•β–ˆβ–ˆβ•”β•β•β•β–ˆβ–ˆβ•”β•β•β•β•β•β–ˆβ–ˆβ•‘ β–ˆβ–ˆβ•‘ β–ˆβ–ˆβ–ˆβ–ˆβ•— β–ˆβ–ˆβ–ˆβ–ˆβ•‘ - β–ˆβ–ˆβ•‘ β–ˆβ–ˆβ•‘ β–ˆβ–ˆβ•‘ β–ˆβ–ˆβ–ˆβ–ˆβ–ˆβ•— β–ˆβ–ˆβ•‘ β–ˆβ–ˆβ•‘ β–ˆβ–ˆβ•”β–ˆβ–ˆβ–ˆβ–ˆβ•”β–ˆβ–ˆβ•‘ - β–ˆβ–ˆβ•‘ β–ˆβ–ˆβ•‘ β–ˆβ–ˆβ•‘ β–ˆβ–ˆβ•”β•β•β• β–ˆβ–ˆβ•‘ β–ˆβ–ˆβ•‘ β–ˆβ–ˆβ•‘β•šβ–ˆβ–ˆβ•”β•β–ˆβ–ˆβ•‘ - β–ˆβ–ˆβ–ˆβ–ˆβ–ˆβ–ˆβ–ˆβ•—β–ˆβ–ˆβ•‘ β–ˆβ–ˆβ•‘ β–ˆβ–ˆβ–ˆβ–ˆβ–ˆβ–ˆβ–ˆβ•—β–ˆβ–ˆβ–ˆβ–ˆβ–ˆβ–ˆβ–ˆβ•—β–ˆβ–ˆβ–ˆβ–ˆβ–ˆβ–ˆβ–ˆβ•—β–ˆβ–ˆβ•‘ β•šβ•β• β–ˆβ–ˆβ•‘ - β•šβ•β•β•β•β•β•β•β•šβ•β• β•šβ•β• β•šβ•β•β•β•β•β•β•β•šβ•β•β•β•β•β•β•β•šβ•β•β•β•β•β•β•β•šβ•β• β•šβ•β• - -17:27:32 - LiteLLM:DEBUG: utils.py:5620 - checking potential_model_names in litellm.model_cost: {'split_model': '981a6c3e187a3f563db301b17f0776519148f13a3a80f13658446ed9332f8f22', 'combined_model_name': '981a6c3e187a3f563db301b17f0776519148f13a3a80f13658446ed9332f8f22', 'stripped_model_name': '981a6c3e187a3f563db301b17f0776519148f13a3a80f13658446ed9332f8f22', 'combined_stripped_model_name': '981a6c3e187a3f563db301b17f0776519148f13a3a80f13658446ed9332f8f22', 'custom_llm_provider': None} -17:27:32 - LiteLLM:DEBUG: utils.py:5915 - Error getting model info: This model isn't mapped yet. Add it here - https://github.com/BerriAI/litellm/blob/main/model_prices_and_context_window.json -17:27:32 - LiteLLM:DEBUG: utils.py:2900 - added/updated model=981a6c3e187a3f563db301b17f0776519148f13a3a80f13658446ed9332f8f22 in litellm.model_cost: 981a6c3e187a3f563db301b17f0776519148f13a3a80f13658446ed9332f8f22 -17:27:32 - LiteLLM:DEBUG: utils.py:5620 - checking potential_model_names in litellm.model_cost: {'split_model': 'oia-gpt-realtime', 'combined_model_name': 'azure/oia-gpt-realtime', 'stripped_model_name': 'azure/oia-gpt-realtime', 'combined_stripped_model_name': 'azure/oia-gpt-realtime', 'custom_llm_provider': 'azure'} -17:27:32 - LiteLLM:DEBUG: utils.py:5915 - Error getting model info: This model isn't mapped yet. Add it here - https://github.com/BerriAI/litellm/blob/main/model_prices_and_context_window.json -17:27:32 - LiteLLM:DEBUG: utils.py:2900 - added/updated model=azure/oia-gpt-realtime in litellm.model_cost: azure/oia-gpt-realtime -17:27:32 - LiteLLM:DEBUG: utils.py:5620 - checking potential_model_names in litellm.model_cost: {'split_model': '91a8591ffdee6edec7016031dc8a092b14ce59f85b1b105c0afa69280775d380', 'combined_model_name': '91a8591ffdee6edec7016031dc8a092b14ce59f85b1b105c0afa69280775d380', 'stripped_model_name': '91a8591ffdee6edec7016031dc8a092b14ce59f85b1b105c0afa69280775d380', 'combined_stripped_model_name': '91a8591ffdee6edec7016031dc8a092b14ce59f85b1b105c0afa69280775d380', 'custom_llm_provider': None} -17:27:32 - LiteLLM:DEBUG: utils.py:5915 - Error getting model info: This model isn't mapped yet. Add it here - https://github.com/BerriAI/litellm/blob/main/model_prices_and_context_window.json -17:27:32 - LiteLLM:DEBUG: utils.py:2900 - added/updated model=91a8591ffdee6edec7016031dc8a092b14ce59f85b1b105c0afa69280775d380 in litellm.model_cost: 91a8591ffdee6edec7016031dc8a092b14ce59f85b1b105c0afa69280775d380 -17:27:32 - LiteLLM:DEBUG: utils.py:5620 - checking potential_model_names in litellm.model_cost: {'split_model': 'claude-opus-4-7', 'combined_model_name': 'anthropic/claude-opus-4-7', 'stripped_model_name': 'anthropic/claude-opus-4-7', 'combined_stripped_model_name': 'anthropic/claude-opus-4-7', 'custom_llm_provider': 'anthropic'} -17:27:32 - LiteLLM:DEBUG: utils.py:2900 - added/updated model=claude-opus-4-7 in litellm.model_cost: claude-opus-4-7 -17:27:32 - LiteLLM Router:DEBUG: router.py:7008 - -Initialized Model List ['gpt-realtime', 'claude-opus-4-7'] -17:27:32 - LiteLLM Router:INFO: router.py:802 - Routing strategy: simple-shuffle -17:27:32 - LiteLLM Proxy:DEBUG: proxy_server.py:3838 - Policy engine: no policies in config, skipping -17:27:32 - LiteLLM Proxy:DEBUG: utils.py:2455 - Creating Prisma Client.. -17:27:32 - LiteLLM Proxy:DEBUG: utils.py:2522 - Success - Created Prisma Client -17:27:32 - LiteLLM Proxy:DEBUG: utils.py:3740 - PrismaClient: connect() called Attempting to Connect to DB -17:27:32 - LiteLLM Proxy:DEBUG: utils.py:3744 - PrismaClient: DB not connected, Attempting to Connect to DB -query-engine ac9d7041ed77bcc8a8dbd2ab6616b39013829574 -17:27:32 - LiteLLM Proxy:DEBUG: prisma_client.py:247 - IAM token auth not enabled, skipping token refresh task -17:27:32 - LiteLLM Proxy:INFO: utils.py:4302 - Started Prisma DB health watchdog (interval=30s, reconnect_cooldown=15s, probe_timeout=5.0s, reconnect_timeout=30.0s) -17:27:32 - LiteLLM Proxy:INFO: utils.py:4057 - Found prisma-query-engine at PID 66228. -17:27:32 - LiteLLM Proxy:INFO: utils.py:4061 - Watching engine PID 66228 via waitpid thread. -17:27:32 - LiteLLM:DEBUG: logging_callback_manager.py:336 - Custom logger of type SkillsInjectionHook, key: SkillsInjectionHook-max_iterations=10-sandbox_timeout=120-message_logging=True-turn_off_message_logging=False already exists in [, , , , , , ], not adding again.. -17:27:32 - LiteLLM Proxy:DEBUG: proxy_server.py:896 - About to initialize semantic tool filter -17:27:32 - LiteLLM Proxy:DEBUG: proxy_server.py:899 - litellm_settings keys = [] -17:27:32 - LiteLLM Proxy:DEBUG: proxy_server.py:6143 - Semantic tool filter not configured or not enabled, skipping initialization -17:27:32 - LiteLLM Proxy:DEBUG: proxy_server.py:906 - After semantic tool filter initialization -17:27:32 - LiteLLM Proxy:DEBUG: proxy_server.py:920 - prisma_client: -17:27:32 - LiteLLM Proxy:INFO: proxy_server.py:6390 - Tag spend update job scheduled at 34s interval (2.3x main job interval) -17:27:32 - LiteLLM Proxy:DEBUG: hanging_request_check.py:148 - Checking for hanging requests.... -17:27:32 - LiteLLM Proxy:INFO: utils.py:5078 - Starting spend logs queue monitor (threshold: 100, poll_interval: 2.0s) -17:27:32 - LiteLLM Proxy:INFO: utils.py:2629 - All necessary views exist! -17:27:32 - LiteLLM Proxy:INFO: proxy_server.py:875 - Password migration: No plaintext passwords found -17:27:32 - LiteLLM Proxy:DEBUG: proxy_server.py:4179 - len new_models: 0 -17:27:32 - LiteLLM Proxy:DEBUG: proxy_server.py:5384 - guardrails from the DB [] -17:27:32 - LiteLLM Proxy:INFO: policy_registry.py:577 - Synced 0 production policies and 0 draft/published (by ID) from DB to in-memory registry -17:27:32 - LiteLLM Proxy:INFO: attachment_registry.py:481 - Synced 0 attachments from DB to in-memory registry -17:27:32 - LiteLLM Proxy:DEBUG: proxy_server.py:5420 - Successfully synced policies and attachments from DB -17:27:32 - LiteLLM:DEBUG: mcp_server_manager.py:2644 - Loading MCP servers from database into registry... -17:27:32 - LiteLLM:INFO: mcp_server_manager.py:2664 - Found 0 MCP servers in database -17:27:32 - LiteLLM:DEBUG: mcp_server_manager.py:2697 - MCP registry refreshed (0 servers in registry) -17:27:32 - LiteLLM Proxy:DEBUG: pass_through_endpoints.py:2289 - initializing pass through endpoints -17:27:32 - LiteLLM Proxy:DEBUG: prompt_registry.py:54 - Discovering prompt integrations in: litellm.integrations.weave -17:27:32 - LiteLLM Proxy:DEBUG: prompt_registry.py:54 - Discovering prompt integrations in: litellm.integrations.litellm_agent -17:27:32 - LiteLLM Proxy:DEBUG: prompt_registry.py:54 - Discovering prompt integrations in: litellm.integrations.dotprompt -17:27:32 - LiteLLM Proxy:DEBUG: prompt_registry.py:65 - Found prompt_initializer_registry in litellm.integrations.dotprompt: ['dotprompt'] -17:27:32 - LiteLLM Proxy:DEBUG: prompt_registry.py:54 - Discovering prompt integrations in: litellm.integrations.gitlab -17:27:32 - LiteLLM Proxy:DEBUG: prompt_registry.py:65 - Found prompt_initializer_registry in litellm.integrations.gitlab: ['gitlab'] -17:27:32 - LiteLLM Proxy:DEBUG: prompt_registry.py:54 - Discovering prompt integrations in: litellm.integrations.azure_sentinel -17:27:32 - LiteLLM Proxy:DEBUG: prompt_registry.py:54 - Discovering prompt integrations in: litellm.integrations.arize -17:27:32 - LiteLLM Proxy:DEBUG: prompt_registry.py:65 - Found prompt_initializer_registry in litellm.integrations.arize: ['arize_phoenix'] -17:27:32 - LiteLLM Proxy:DEBUG: prompt_registry.py:54 - Discovering prompt integrations in: litellm.integrations.agentops -17:27:32 - LiteLLM Proxy:DEBUG: prompt_registry.py:54 - Discovering prompt integrations in: litellm.integrations.focus -17:27:32 - LiteLLM Proxy:DEBUG: prompt_registry.py:54 - Discovering prompt integrations in: litellm.integrations.prometheus_helpers -17:27:32 - LiteLLM Proxy:DEBUG: prompt_registry.py:54 - Discovering prompt integrations in: litellm.integrations.generic_prompt_management -17:27:32 - LiteLLM Proxy:DEBUG: prompt_registry.py:65 - Found prompt_initializer_registry in litellm.integrations.generic_prompt_management: ['generic_prompt_management'] -17:27:32 - LiteLLM Proxy:DEBUG: prompt_registry.py:54 - Discovering prompt integrations in: litellm.integrations.levo -17:27:32 - LiteLLM Proxy:DEBUG: prompt_registry.py:54 - Discovering prompt integrations in: litellm.integrations.websearch_interception -17:27:32 - LiteLLM Proxy:DEBUG: prompt_registry.py:54 - Discovering prompt integrations in: litellm.integrations.deepeval -17:27:32 - LiteLLM Proxy:DEBUG: prompt_registry.py:54 - Discovering prompt integrations in: litellm.integrations.bitbucket -17:27:32 - LiteLLM Proxy:DEBUG: prompt_registry.py:65 - Found prompt_initializer_registry in litellm.integrations.bitbucket: ['bitbucket'] -17:27:32 - LiteLLM Proxy:DEBUG: prompt_registry.py:54 - Discovering prompt integrations in: litellm.integrations.vantage -17:27:32 - LiteLLM Proxy:DEBUG: prompt_registry.py:76 - Discovered 5 prompt initializers: ['dotprompt', 'gitlab', 'arize_phoenix', 'generic_prompt_management', 'bitbucket'] -17:27:32 - LiteLLM Proxy:INFO: proxy_server.py:5563 - Loading 0 search tool(s) from database into router -17:27:32 - LiteLLM Proxy:DEBUG: proxy_server.py:5583 - No search tools found in database, keeping config-loaded search tools (if any) -17:27:32 - LiteLLM Proxy:INFO: tool_registry_writer.py:329 - ToolPolicyRegistry: synced 15 tool policies and 0 object permissions from DB -17:27:32 - LiteLLM Proxy:DEBUG: proxy_server.py:5440 - Successfully synced tool policy from DB -17:27:32 - LiteLLM:DEBUG: focus_logger.py:167 - No Focus export logger registered; skipping scheduler -17:27:32 - LiteLLM Proxy:DEBUG: proxy_server.py:6678 - key_rotation_enabled: False -17:27:32 - LiteLLM Proxy:DEBUG: proxy_server.py:6714 - Key rotation disabled (set LITELLM_KEY_ROTATION_ENABLED=true to enable) -17:27:32 - LiteLLM Proxy:INFO: proxy_server.py:6545 - Batch cost check job scheduled successfully -17:27:32 - LiteLLM Proxy:INFO: proxy_server.py:6576 - Responses cost check job scheduled successfully -17:27:32 - LiteLLM Proxy:INFO: proxy_server.py:6595 - APScheduler started with memory leak prevention settings: removed jitter, increased intervals, misfire_grace_time=3600 -17:27:32 - LiteLLM Proxy:DEBUG: proxy_server.py:6880 - LiteLLM: Pyroscope profiling is disabled (set LITELLM_ENABLE_PYROSCOPE=true to enable). -17:27:32 - LiteLLM Proxy:INFO: proxy_server.py:756 - SESSION REUSE: Created shared aiohttp session for connection pooling (ID: 4796655632, limit=1000, limit_per_host=500) -INFO: Application startup complete. -INFO: Uvicorn running on http://0.0.0.0:4000 (Press CTRL+C to quit) -17:27:39 - LiteLLM Proxy:DEBUG: http_parsing_utils.py:529 - populate_request_with_path_params: No vector_store_id present in path=/v1/chat/completions -17:27:39 - LiteLLM:DEBUG: user_api_key_auth.py:1032 - api key not found in cache. -17:27:39 - LiteLLM Proxy:DEBUG: common_request_processing.py:876 - Request received by LiteLLM: -{ - "model": "claude-opus-4-7", - "messages": [ - { - "role": "user", - "content": "Hello, how are you?" - } - ], - "reasoning_effort": "max", - "metadata": { - "user_api_key_user_id": "default_user_id" - } -} -17:27:39 - LiteLLM Proxy:DEBUG: litellm_pre_call_utils.py:971 - Request Headers: {'content-type': 'application/json', 'user-agent': 'PostmanRuntime/7.53.0', 'accept': '*/*', 'postman-token': '08d81a08-addf-4a27-a9f8-302f317d7296', 'host': '0.0.0.0:4000', 'accept-encoding': 'gzip, deflate, br', 'connection': 'keep-alive', 'content-length': '183'} -17:27:39 - LiteLLM Proxy:DEBUG: litellm_pre_call_utils.py:972 - Raw Headers: RedactedDict(REDACTED) -17:27:39 - LiteLLM Proxy:DEBUG: litellm_pre_call_utils.py:1065 - receiving data: {'model': 'claude-opus-4-7', 'messages': [{'role': 'user', 'content': 'Hello, how are you?'}], 'reasoning_effort': 'max', 'metadata': {'headers': {'content-type': 'application/json', 'user-agent': 'PostmanRuntime/7.53.0', 'accept': '*/*', 'postman-token': '08d81a08-addf-4a27-a9f8-302f317d7296', 'host': '0.0.0.0:4000', 'accept-encoding': 'gzip, deflate, br', 'connection': 'keep-alive', 'content-length': '183'}}, 'proxy_server_request': {'url': 'http://0.0.0.0:4000/v1/chat/completions', 'method': 'POST', 'headers': {'content-type': 'application/json', 'user-agent': 'PostmanRuntime/7.53.0', 'accept': '*/*', 'postman-token': '08d81a08-addf-4a27-a9f8-302f317d7296', 'host': '0.0.0.0:4000', 'accept-encoding': 'gzip, deflate, br', 'connection': 'keep-alive', 'content-length': '183'}, 'body': None, 'arrival_time': 1776686259.683504}, 'secret_fields': {'raw_headers': RedactedDict(REDACTED)}} -17:27:39 - LiteLLM Proxy:DEBUG: litellm_pre_call_utils.py:2246 - Policy engine: registry initialized=True, policy_count=0 -17:27:39 - LiteLLM Proxy:DEBUG: litellm_pre_call_utils.py:2267 - Policy engine: matching policies for context team_alias=None, key_alias=None, model=claude-opus-4-7, tags=None -17:27:39 - LiteLLM Proxy:DEBUG: litellm_pre_call_utils.py:2104 - Policy engine: matched policies via attachments: [] -17:27:39 - LiteLLM Proxy:DEBUG: policy_resolver.py:173 - No policies match context: team_alias=None, key_alias=None, model=claude-opus-4-7 -17:27:39 - LiteLLM Proxy:DEBUG: litellm_pre_call_utils.py:2165 - Policy engine: resolved guardrails: [] -17:27:39 - LiteLLM Proxy:DEBUG: litellm_pre_call_utils.py:1398 - [PROXY] returned data from litellm_pre_call_utils: {'model': 'claude-opus-4-7', 'messages': [{'role': 'user', 'content': 'Hello, how are you?'}], 'reasoning_effort': 'max', 'metadata': {'headers': {'content-type': 'application/json', 'user-agent': 'PostmanRuntime/7.53.0', 'accept': '*/*', 'postman-token': '08d81a08-addf-4a27-a9f8-302f317d7296', 'host': '0.0.0.0:4000', 'accept-encoding': 'gzip, deflate, br', 'connection': 'keep-alive', 'content-length': '183'}, 'requester_metadata': {'headers': {'content-type': 'application/json', 'user-agent': 'PostmanRuntime/7.53.0', 'accept': '*/*', 'postman-token': '08d81a08-addf-4a27-a9f8-302f317d7296', 'host': '0.0.0.0:4000', 'accept-encoding': 'gzip, deflate, br', 'connection': 'keep-alive', 'content-length': '183'}}, 'user_api_key_hash': '88dc28d0f030c55ed4ab77ed8faf098196cb1c05df778539800c9f1243fe6b4b', 'user_api_key_alias': None, 'user_api_key_spend': 0.0, 'user_api_key_max_budget': None, 'user_api_key_team_id': None, 'user_api_key_project_id': None, 'user_api_key_project_alias': None, 'user_api_key_user_id': 'default_user_id', 'user_api_key_org_id': None, 'user_api_key_org_alias': None, 'user_api_key_team_alias': None, 'user_api_key_end_user_id': None, 'user_api_key_user_email': None, 'user_api_key_request_route': '/v1/chat/completions', 'user_api_key_budget_reset_at': None, 'user_api_key_auth_metadata': {}, 'user_REDACTED', 'agent_id': None, 'user_api_end_user_max_budget': None, 'user_api_key_auth': UserAPIKeyAuth(token='88dc28d0f030c55ed4ab77ed8faf098196cb1c05df778539800c9f1243fe6b4b', key_name=None, key_alias=None, spend=0.0, max_budget=None, expires=None, models=[], aliases={}, config={}, user_id='default_user_id', team_id=None, agent_id=None, project_id=None, max_parallel_requests=None, metadata={}, tpm_limit=None, rpm_limit=None, budget_duration=None, budget_reset_at=None, allowed_cache_controls=[], allowed_routes=[], permissions={}, model_spend={}, model_max_budget={}, soft_budget_cooldown=False, blocked=None, litellm_budget_table=None, org_id=None, created_at=None, created_by=None, updated_at=None, updated_by=None, last_active=None, object_permission_id=None, object_permission=None, access_group_ids=None, rotation_count=0, auto_rotate=False, rotation_interval=None, last_rotation_at=None, key_rotation_at=None, router_settings=None, budget_limits=None, team_spend=None, team_alias=None, team_tpm_limit=None, team_rpm_limit=None, team_max_budget=None, team_soft_budget=None, team_models=[], team_blocked=False, soft_budget=None, team_model_aliases=None, team_member=None, team_metadata=None, team_object_permission_id=None, team_member_spend=None, team_member_tpm_limit=None, team_member_rpm_limit=None, end_user_id=None, end_user_tpm_limit=None, end_user_rpm_limit=None, end_user_max_budget=None, end_user_model_max_budget=None, organization_alias=None, organization_max_budget=None, organization_tpm_limit=None, organization_rpm_limit=None, organization_metadata=None, project_alias=None, project_metadata=None, last_refreshed_at=None, REDACTED', user_role=, allowed_model_region=None, parent_otel_span=None, rpm_limit_per_model=None, tpm_limit_per_model=None, user_tpm_limit=None, user_rpm_limit=None, user_email=None, user_spend=None, user_max_budget=None, request_route='/v1/chat/completions', user=None, created_by_user=None, end_user_object_permission=None, jwt_claims=None), 'litellm_api_version': '1.83.9', 'global_max_parallel_requests': None, 'user_api_key_team_max_budget': None, 'user_api_key_team_spend': None, 'user_api_key_model_max_budget': {}, 'user_api_key_end_user_model_max_budget': None, 'user_api_key_user_spend': None, 'user_api_key_user_max_budget': None, 'user_api_key_metadata': {}, 'user_api_key_team_metadata': None, 'user_api_key_object_permission_id': None, 'user_api_key_team_object_permission_id': None, 'endpoint': 'http://0.0.0.0:4000/v1/chat/completions', 'litellm_parent_otel_span': None, 'requester_ip_address': '127.0.0.1', 'user_agent': 'PostmanRuntime/7.53.0'}, 'proxy_server_request': {'url': 'http://0.0.0.0:4000/v1/chat/completions', 'method': 'POST', 'headers': {'content-type': 'application/json', 'user-agent': 'PostmanRuntime/7.53.0', 'accept': '*/*', 'postman-token': '08d81a08-addf-4a27-a9f8-302f317d7296', 'host': '0.0.0.0:4000', 'accept-encoding': 'gzip, deflate, br', 'connection': 'keep-alive', 'content-length': '183'}, 'body': {'model': 'claude-opus-4-7', 'messages': [{'role': 'user', 'content': 'Hello, how are you?'}], 'reasoning_effort': 'max', 'metadata': {'headers': {'content-type': 'application/json', 'user-agent': 'PostmanRuntime/7.53.0', 'accept': '*/*', 'postman-token': '08d81a08-addf-4a27-a9f8-302f317d7296', 'host': '0.0.0.0:4000', 'accept-encoding': 'gzip, deflate, br', 'connection': 'keep-alive', 'content-length': '183'}, 'requester_metadata': {'headers': {'content-type': 'application/json', 'user-agent': 'PostmanRuntime/7.53.0', 'accept': '*/*', 'postman-token': '08d81a08-addf-4a27-a9f8-302f317d7296', 'host': '0.0.0.0:4000', 'accept-encoding': 'gzip, deflate, br', 'connection': 'keep-alive', 'content-length': '183'}}, 'user_api_key_hash': '88dc28d0f030c55ed4ab77ed8faf098196cb1c05df778539800c9f1243fe6b4b', 'user_api_key_alias': None, 'user_api_key_spend': 0.0, 'user_api_key_max_budget': None, 'user_api_key_team_id': None, 'user_api_key_project_id': None, 'user_api_key_project_alias': None, 'user_api_key_user_id': 'default_user_id', 'user_api_key_org_id': None, 'user_api_key_org_alias': None, 'user_api_key_team_alias': None, 'user_api_key_end_user_id': None, 'user_api_key_user_email': None, 'user_api_key_request_route': '/v1/chat/completions', 'user_api_key_budget_reset_at': None, 'user_api_key_auth_metadata': {}, 'user_REDACTED', 'agent_id': None, 'user_api_end_user_max_budget': None, 'user_api_key_auth': UserAPIKeyAuth(token='88dc28d0f030c55ed4ab77ed8faf098196cb1c05df778539800c9f1243fe6b4b', key_name=None, key_alias=None, spend=0.0, max_budget=None, expires=None, models=[], aliases={}, config={}, user_id='default_user_id', team_id=None, agent_id=None, project_id=None, max_parallel_requests=None, metadata={}, tpm_limit=None, rpm_limit=None, budget_duration=None, budget_reset_at=None, allowed_cache_controls=[], allowed_routes=[], permissions={}, model_spend={}, model_max_budget={}, soft_budget_cooldown=False, blocked=None, litellm_budget_table=None, org_id=None, created_at=None, created_by=None, updated_at=None, updated_by=None, last_active=None, object_permission_id=None, object_permission=None, access_group_ids=None, rotation_count=0, auto_rotate=False, rotation_interval=None, last_rotation_at=None, key_rotation_at=None, router_settings=None, budget_limits=None, team_spend=None, team_alias=None, team_tpm_limit=None, team_rpm_limit=None, team_max_budget=None, team_soft_budget=None, team_models=[], team_blocked=False, soft_budget=None, team_model_aliases=None, team_member=None, team_metadata=None, team_object_permission_id=None, team_member_spend=None, team_member_tpm_limit=None, team_member_rpm_limit=None, end_user_id=None, end_user_tpm_limit=None, end_user_rpm_limit=None, end_user_max_budget=None, end_user_model_max_budget=None, organization_alias=None, organization_max_budget=None, organization_tpm_limit=None, organization_rpm_limit=None, organization_metadata=None, project_alias=None, project_metadata=None, last_refreshed_at=None, REDACTED', user_role=, allowed_model_region=None, parent_otel_span=None, rpm_limit_per_model=None, tpm_limit_per_model=None, user_tpm_limit=None, user_rpm_limit=None, user_email=None, user_spend=None, user_max_budget=None, request_route='/v1/chat/completions', user=None, created_by_user=None, end_user_object_permission=None, jwt_claims=None), 'litellm_api_version': '1.83.9', 'global_max_parallel_requests': None, 'user_api_key_team_max_budget': None, 'user_api_key_team_spend': None, 'user_api_key_model_max_budget': {}, 'user_api_key_end_user_model_max_budget': None, 'user_api_key_user_spend': None, 'user_api_key_user_max_budget': None, 'user_api_key_metadata': {}, 'user_api_key_team_metadata': None, 'user_api_key_object_permission_id': None, 'user_api_key_team_object_permission_id': None, 'endpoint': 'http://0.0.0.0:4000/v1/chat/completions', 'litellm_parent_otel_span': None, 'requester_ip_address': '127.0.0.1', 'user_agent': 'PostmanRuntime/7.53.0'}, 'proxy_server_request': {...}, 'secret_fields': {'raw_headers': RedactedDict(REDACTED)}}, 'arrival_time': 1776686259.683504}, 'secret_fields': {'raw_headers': RedactedDict(REDACTED)}} -17:27:39 - LiteLLM:DEBUG: logging_callback_manager.py:336 - Custom logger of type _ProxyDBLogger, key: _ProxyDBLogger-message_logging=True-turn_off_message_logging=False already exists in [>, , , ], not adding again.. -17:27:39 - LiteLLM:DEBUG: utils.py:482 - Initialized litellm callbacks, Async Success Callbacks: [>, , , , , , , , , , , , ] -17:27:39 - LiteLLM:DEBUG: utils.py:1008 - Removing thought signatures from tool call IDs for non-Gemini model -17:27:39 - LiteLLM:DEBUG: litellm_logging.py:557 - self.optional_params: {} -17:27:39 - LiteLLM Proxy:DEBUG: utils.py:1346 - Inside Proxy Logging Pre-call hook! -17:27:39 - LiteLLM Proxy:DEBUG: max_budget_limiter.py:23 - Inside Max Budget Limiter Pre-Call Hook -17:27:39 - LiteLLM Proxy:DEBUG: parallel_request_limiter_v3.py:1288 - Inside Rate Limit Pre-Call Hook -17:27:39 - LiteLLM Proxy:DEBUG: cache_control_check.py:27 - Inside Cache Control Check Pre-Call Hook -17:27:39 - LiteLLM Proxy:INFO: route_llm_request.py:164 - SESSION REUSE: Attached shared aiohttp session to request (ID: 4796655632) -17:27:39 - LiteLLM Router:DEBUG: router.py:5679 - Inside async function with retries. -17:27:39 - LiteLLM Router:DEBUG: router.py:5703 - async function w/ retries: original_function - >, num_retries - 2 -17:27:39 - LiteLLM Router:DEBUG: router.py:9140 - initial list of deployments: [{'model_name': 'claude-opus-4-7', 'litellm_params': {'use_in_pass_through': False, 'use_litellm_proxy': False, 'merge_reasoning_content_in_choices': False, 'model': 'anthropic/claude-opus-4-7'}, 'model_info': {'id': '91a8591ffdee6edec7016031dc8a092b14ce59f85b1b105c0afa69280775d380', 'db_model': False}}] -17:27:39 - LiteLLM Router:DEBUG: router.py:9216 - healthy_deployments after team filter: [{'model_name': 'claude-opus-4-7', 'litellm_params': {'use_in_pass_through': False, 'use_litellm_proxy': False, 'merge_reasoning_content_in_choices': False, 'model': 'anthropic/claude-opus-4-7'}, 'model_info': {'id': '91a8591ffdee6edec7016031dc8a092b14ce59f85b1b105c0afa69280775d380', 'db_model': False}}] -17:27:39 - LiteLLM Router:DEBUG: router.py:9226 - healthy_deployments after web search filter: [{'model_name': 'claude-opus-4-7', 'litellm_params': {'use_in_pass_through': False, 'use_litellm_proxy': False, 'merge_reasoning_content_in_choices': False, 'model': 'anthropic/claude-opus-4-7'}, 'model_info': {'id': '91a8591ffdee6edec7016031dc8a092b14ce59f85b1b105c0afa69280775d380', 'db_model': False}}] -17:27:39 - LiteLLM Router:DEBUG: cooldown_handlers.py:347 - retrieve cooldown models: [] -17:27:39 - LiteLLM Router:DEBUG: router.py:9245 - cooldown deployments: [] -17:27:39 - LiteLLM Router:DEBUG: router.py:9987 - cooldown deployments: [] -17:27:39 - LiteLLM:DEBUG: utils.py:482 - - -17:27:39 - LiteLLM:DEBUG: utils.py:482 - Request to litellm: -17:27:39 - LiteLLM:DEBUG: utils.py:482 - litellm.acompletion(use_in_pass_through=False, use_litellm_proxy=False, merge_reasoning_content_in_choices=False, model='anthropic/claude-opus-4-7', messages=[{'role': 'user', 'content': 'Hello, how are you?'}], caching=False, client=None, reasoning_effort='max', metadata={'headers': {'content-type': 'application/json', 'user-agent': 'PostmanRuntime/7.53.0', 'accept': '*/*', 'postman-token': '08d81a08-addf-4a27-a9f8-302f317d7296', 'host': '0.0.0.0:4000', 'accept-encoding': 'gzip, deflate, br', 'connection': 'keep-alive', 'content-length': '183'}, 'requester_metadata': {'headers': {'content-type': 'application/json', 'user-agent': 'PostmanRuntime/7.53.0', 'accept': '*/*', 'postman-token': '08d81a08-addf-4a27-a9f8-302f317d7296', 'host': '0.0.0.0:4000', 'accept-encoding': 'gzip, deflate, br', 'connection': 'keep-alive', 'content-length': '183'}}, 'user_api_key_hash': '88dc28d0f030c55ed4ab77ed8faf098196cb1c05df778539800c9f1243fe6b4b', 'user_api_key_alias': None, 'user_api_key_spend': 0.0, 'user_api_key_max_budget': None, 'user_api_key_team_id': None, 'user_api_key_project_id': None, 'user_api_key_project_alias': None, 'user_api_key_user_id': 'default_user_id', 'user_api_key_org_id': None, 'user_api_key_org_alias': None, 'user_api_key_team_alias': None, 'user_api_key_end_user_id': None, 'user_api_key_user_email': None, 'user_api_key_request_route': '/v1/chat/completions', 'user_api_key_budget_reset_at': None, 'user_api_key_auth_metadata': {}, 'user_REDACTED', 'agent_id': None, 'user_api_end_user_max_budget': None, 'user_api_key_auth': UserAPIKeyAuth(token='88dc28d0f030c55ed4ab77ed8faf098196cb1c05df778539800c9f1243fe6b4b', key_name=None, key_alias=None, spend=0.0, max_budget=None, expires=None, models=[], aliases={}, config={}, user_id='default_user_id', team_id=None, agent_id=None, project_id=None, max_parallel_requests=None, metadata={}, tpm_limit=None, rpm_limit=None, budget_duration=None, budget_reset_at=None, allowed_cache_controls=[], allowed_routes=[], permissions={}, model_spend={}, model_max_budget={}, soft_budget_cooldown=False, blocked=None, litellm_budget_table=None, org_id=None, created_at=None, created_by=None, updated_at=None, updated_by=None, last_active=None, object_permission_id=None, object_permission=None, access_group_ids=None, rotation_count=0, auto_rotate=False, rotation_interval=None, last_rotation_at=None, key_rotation_at=None, router_settings=None, budget_limits=None, team_spend=None, team_alias=None, team_tpm_limit=None, team_rpm_limit=None, team_max_budget=None, team_soft_budget=None, team_models=[], team_blocked=False, soft_budget=None, team_model_aliases=None, team_member=None, team_metadata=None, team_object_permission_id=None, team_member_spend=None, team_member_tpm_limit=None, team_member_rpm_limit=None, end_user_id=None, end_user_tpm_limit=None, end_user_rpm_limit=None, end_user_max_budget=None, end_user_model_max_budget=None, organization_alias=None, organization_max_budget=None, organization_tpm_limit=None, organization_rpm_limit=None, organization_metadata=None, project_alias=None, project_metadata=None, last_refreshed_at=1776686259.698418, REDACTED', user_role=, allowed_model_region=None, parent_otel_span=None, rpm_limit_per_model=None, tpm_limit_per_model=None, user_tpm_limit=None, user_rpm_limit=None, user_email=None, user_spend=None, user_max_budget=None, request_route='/v1/chat/completions', user=None, created_by_user=None, end_user_object_permission=None, jwt_claims=None), 'litellm_api_version': '1.83.9', 'global_max_parallel_requests': None, 'user_api_key_team_max_budget': None, 'user_api_key_team_spend': None, 'user_api_key_model_max_budget': {}, 'user_api_key_end_user_model_max_budget': None, 'user_api_key_user_spend': None, 'user_api_key_user_max_budget': None, 'user_api_key_metadata': {}, 'user_api_key_team_metadata': None, 'user_api_key_object_permission_id': None, 'user_api_key_team_object_permission_id': None, 'endpoint': 'http://0.0.0.0:4000/v1/chat/completions', 'litellm_parent_otel_span': None, 'requester_ip_address': '127.0.0.1', 'user_agent': 'PostmanRuntime/7.53.0', 'queue_time_seconds': 0.009355783462524414, 'model_group': 'claude-opus-4-7', 'model_group_alias': None, 'model_group_size': 1, 'attempted_retries': 0, 'max_retries': 2, 'deployment': 'anthropic/claude-opus-4-7', 'model_info': {'id': '91a8591ffdee6edec7016031dc8a092b14ce59f85b1b105c0afa69280775d380', 'db_model': False}, 'api_base': None, 'deployment_model_name': 'claude-opus-4-7', 'caching_groups': None}, proxy_server_request={'url': 'http://0.0.0.0:4000/v1/chat/completions', 'method': 'POST', 'headers': {'content-type': 'application/json', 'user-agent': 'PostmanRuntime/7.53.0', 'accept': '*/*', 'postman-token': '08d81a08-addf-4a27-a9f8-302f317d7296', 'host': '0.0.0.0:4000', 'accept-encoding': 'gzip, deflate, br', 'connection': 'keep-alive', 'content-length': '183'}, 'body': {'model': 'claude-opus-4-7', 'messages': [{'role': 'user', 'content': 'Hello, how are you?'}], 'reasoning_effort': 'max', 'metadata': {'headers': {'content-type': 'application/json', 'user-agent': 'PostmanRuntime/7.53.0', 'accept': '*/*', 'postman-token': '08d81a08-addf-4a27-a9f8-302f317d7296', 'host': '0.0.0.0:4000', 'accept-encoding': 'gzip, deflate, br', 'connection': 'keep-alive', 'content-length': '183'}, 'requester_metadata': {'headers': {'content-type': 'application/json', 'user-agent': 'PostmanRuntime/7.53.0', 'accept': '*/*', 'postman-token': '08d81a08-addf-4a27-a9f8-302f317d7296', 'host': '0.0.0.0:4000', 'accept-encoding': 'gzip, deflate, br', 'connection': 'keep-alive', 'content-length': '183'}}, 'user_api_key_hash': '88dc28d0f030c55ed4ab77ed8faf098196cb1c05df778539800c9f1243fe6b4b', 'user_api_key_alias': None, 'user_api_key_spend': 0.0, 'user_api_key_max_budget': None, 'user_api_key_team_id': None, 'user_api_key_project_id': None, 'user_api_key_project_alias': None, 'user_api_key_user_id': 'default_user_id', 'user_api_key_org_id': None, 'user_api_key_org_alias': None, 'user_api_key_team_alias': None, 'user_api_key_end_user_id': None, 'user_api_key_user_email': None, 'user_api_key_request_route': '/v1/chat/completions', 'user_api_key_budget_reset_at': None, 'user_api_key_auth_metadata': {}, 'user_REDACTED', 'agent_id': None, 'user_api_end_user_max_budget': None, 'user_api_key_auth': UserAPIKeyAuth(token='88dc28d0f030c55ed4ab77ed8faf098196cb1c05df778539800c9f1243fe6b4b', key_name=None, key_alias=None, spend=0.0, max_budget=None, expires=None, models=[], aliases={}, config={}, user_id='default_user_id', team_id=None, agent_id=None, project_id=None, max_parallel_requests=None, metadata={}, tpm_limit=None, rpm_limit=None, budget_duration=None, budget_reset_at=None, allowed_cache_controls=[], allowed_routes=[], permissions={}, model_spend={}, model_max_budget={}, soft_budget_cooldown=False, blocked=None, litellm_budget_table=None, org_id=None, created_at=None, created_by=None, updated_at=None, updated_by=None, last_active=None, object_permission_id=None, object_permission=None, access_group_ids=None, rotation_count=0, auto_rotate=False, rotation_interval=None, last_rotation_at=None, key_rotation_at=None, router_settings=None, budget_limits=None, team_spend=None, team_alias=None, team_tpm_limit=None, team_rpm_limit=None, team_max_budget=None, team_soft_budget=None, team_models=[], team_blocked=False, soft_budget=None, team_model_aliases=None, team_member=None, team_metadata=None, team_object_permission_id=None, team_member_spend=None, team_member_tpm_limit=None, team_member_rpm_limit=None, end_user_id=None, end_user_tpm_limit=None, end_user_rpm_limit=None, end_user_max_budget=None, end_user_model_max_budget=None, organization_alias=None, organization_max_budget=None, organization_tpm_limit=None, organization_rpm_limit=None, organization_metadata=None, project_alias=None, project_metadata=None, last_refreshed_at=1776686259.698418, REDACTED', user_role=, allowed_model_region=None, parent_otel_span=None, rpm_limit_per_model=None, tpm_limit_per_model=None, user_tpm_limit=None, user_rpm_limit=None, user_email=None, user_spend=None, user_max_budget=None, request_route='/v1/chat/completions', user=None, created_by_user=None, end_user_object_permission=None, jwt_claims=None), 'litellm_api_version': '1.83.9', 'global_max_parallel_requests': None, 'user_api_key_team_max_budget': None, 'user_api_key_team_spend': None, 'user_api_key_model_max_budget': {}, 'user_api_key_end_user_model_max_budget': None, 'user_api_key_user_spend': None, 'user_api_key_user_max_budget': None, 'user_api_key_metadata': {}, 'user_api_key_team_metadata': None, 'user_api_key_object_permission_id': None, 'user_api_key_team_object_permission_id': None, 'endpoint': 'http://0.0.0.0:4000/v1/chat/completions', 'litellm_parent_otel_span': None, 'requester_ip_address': '127.0.0.1', 'user_agent': 'PostmanRuntime/7.53.0', 'queue_time_seconds': 0.009355783462524414, 'model_group': 'claude-opus-4-7', 'model_group_alias': None, 'model_group_size': 1, 'attempted_retries': 0, 'max_retries': 2, 'deployment': 'anthropic/claude-opus-4-7', 'model_info': {'id': '91a8591ffdee6edec7016031dc8a092b14ce59f85b1b105c0afa69280775d380', 'db_model': False}, 'api_base': None, 'deployment_model_name': 'claude-opus-4-7', 'caching_groups': None}, 'proxy_server_request': {...}, 'secret_fields': {'raw_headers': RedactedDict(REDACTED)}}, 'arrival_time': 1776686259.683504}, secret_fields={'raw_headers': RedactedDict(REDACTED)}, litellm_call_id='9d7c129a-8dcc-4fb1-8a30-4d7117085a4f', litellm_logging_obj=, shared_session=, stream=False, litellm_trace_id='066d410c-af0d-4e32-aca5-e3e657fbaff2', model_info={'id': '91a8591ffdee6edec7016031dc8a092b14ce59f85b1b105c0afa69280775d380', 'db_model': False}, timeout=6000.0, max_retries=0) -17:27:39 - LiteLLM:DEBUG: utils.py:482 - - -17:27:39 - LiteLLM:DEBUG: utils.py:482 - ASYNC kwargs[caching]: False; litellm.cache: None; kwargs.get('cache'): None -17:27:39 - LiteLLM:DEBUG: main.py:523 - πŸ”„ SHARED SESSION: acompletion called with shared_session (ID: 4796655632) -17:27:39 - LiteLLM:DEBUG: utils.py:5620 - checking potential_model_names in litellm.model_cost: {'split_model': 'claude-opus-4-7', 'combined_model_name': 'anthropic/claude-opus-4-7', 'stripped_model_name': 'claude-opus-4-7', 'combined_stripped_model_name': 'anthropic/claude-opus-4-7', 'custom_llm_provider': 'anthropic'} -17:27:39 - LiteLLM:INFO: utils.py:4011 - -LiteLLM completion() model= claude-opus-4-7; provider = anthropic -17:27:39 - LiteLLM:DEBUG: utils.py:4014 - -LiteLLM: Params passed to completion() {'model': 'claude-opus-4-7', 'functions': None, 'function_call': None, 'temperature': None, 'top_p': None, 'n': None, 'stream': False, 'stream_options': None, 'stop': None, 'max_tokens': None, 'max_completion_tokens': None, 'modalities': None, 'prediction': None, 'audio': None, 'presence_penalty': None, 'frequency_penalty': None, 'logit_bias': None, 'user': None, 'custom_llm_provider': 'anthropic', 'response_format': None, 'seed': None, 'tools': None, 'tool_choice': None, 'max_retries': 0, 'logprobs': None, 'top_logprobs': None, 'extra_headers': None, 'api_version': None, 'parallel_tool_calls': None, 'drop_params': None, 'allowed_openai_params': None, 'reasoning_effort': 'max', 'verbosity': None, 'additional_drop_params': None, 'messages': [{'role': 'user', 'content': 'Hello, how are you?'}], 'thinking': None, 'web_search_options': None, 'safety_identifier': None, 'service_tier': None} -17:27:39 - LiteLLM:DEBUG: utils.py:4017 - -LiteLLM: Non-Default params passed to completion() {'stream': False, 'max_retries': 0, 'reasoning_effort': 'max'} -17:27:39 - LiteLLM:DEBUG: utils.py:482 - Final returned optional params: {'thinking': {'type': 'adaptive'}, 'output_config': {'effort': 'max'}} -17:27:39 - LiteLLM:DEBUG: litellm_logging.py:557 - self.optional_params: {'stream': False, 'max_retries': 0, 'reasoning_effort': 'max'} -17:27:39 - LiteLLM:DEBUG: utils.py:5620 - checking potential_model_names in litellm.model_cost: {'split_model': 'claude-opus-4-7', 'combined_model_name': 'anthropic/claude-opus-4-7', 'stripped_model_name': 'claude-opus-4-7', 'combined_stripped_model_name': 'anthropic/claude-opus-4-7', 'custom_llm_provider': 'anthropic'} -17:27:39 - LiteLLM:DEBUG: utils.py:5620 - checking potential_model_names in litellm.model_cost: {'split_model': 'claude-opus-4-7', 'combined_model_name': 'anthropic/claude-opus-4-7', 'stripped_model_name': 'claude-opus-4-7', 'combined_stripped_model_name': 'anthropic/claude-opus-4-7', 'custom_llm_provider': 'anthropic'} -17:27:39 - LiteLLM:DEBUG: litellm_logging.py:1160 -  - -POST Request Sent from LiteLLM: -curl -X POST \ -https://api.anthropic.com/v1/messages \ --H 'anthropic-version: 2023-06-01' -H 'accept: application/json' -H 'content-type: application/json' -H 'REDACTED' \ --d '{'model': 'claude-opus-4-7', 'messages': [{'role': 'user', 'content': [{'type': 'text', 'text': 'Hello, how are you?'}]}], 'thinking': {'type': 'adaptive'}, 'output_config': {'effort': 'max'}, 'max_tokens': 128000}' - - -17:27:39 - LiteLLM:DEBUG: main.py:7275 - _is_function_call: False -17:27:39 - LiteLLM:DEBUG: litellm_logging.py:1233 - RAW RESPONSE: - - - -17:27:39 - LiteLLM:DEBUG: http_handler.py:840 - Using AiohttpTransport... -17:27:39 - LiteLLM:DEBUG: http_handler.py:908 - Creating AiohttpTransport... -17:27:39 - LiteLLM:DEBUG: http_handler.py:922 - NEW SESSION: Creating new ClientSession (no shared session provided) -17:27:42 - LiteLLM:DEBUG: litellm_logging.py:1233 - RAW RESPONSE: -{"model":"claude-opus-4-7","id":"msg_018BKbY3htb4ZV1AMbYuQiEe","type":"message","role":"assistant","content":[{"type":"text","text":"Hello! I'm doing well, thanks for asking. I'm ready to help with whatever you needβ€”whether that's answering questions, brainstorming ideas, working through a problem, or just having a conversation. What can I do for you today?"}],"stop_reason":"end_turn","stop_sequence":null,"stop_details":null,"usage":{"input_tokens":19,"cache_creation_input_tokens":0,"cache_read_input_tokens":0,"cache_creation":{"ephemeral_5m_input_tokens":0,"ephemeral_1h_input_tokens":0},"output_tokens":76,"service_tier":"standard","inference_geo":"global"}} - - -17:27:42 - LiteLLM:DEBUG: litellm_logging.py:3312 - Filtered callbacks: [] -17:27:42 - LiteLLM:DEBUG: cost_calculator.py:1134 - selected model name for cost calculation: anthropic/claude-opus-4-7 -17:27:42 - LiteLLM:DEBUG: utils.py:5620 - checking potential_model_names in litellm.model_cost: {'split_model': 'claude-opus-4-7', 'combined_model_name': 'anthropic/claude-opus-4-7', 'stripped_model_name': 'claude-opus-4-7', 'combined_stripped_model_name': 'anthropic/claude-opus-4-7', 'custom_llm_provider': 'anthropic'} -17:27:42 - LiteLLM:DEBUG: litellm_logging.py:1556 - response_cost: 0.001995 -17:27:42 - LiteLLM Router:INFO: router.py:2202 - litellm.acompletion(model=anthropic/claude-opus-4-7) 200 OK -17:27:42 - LiteLLM Router:DEBUG: router.py:5612 - Async Response: ModelResponse(id='chatcmpl-db0b1b06-c051-4281-98ff-6add4a98effb', created=1776686262, model='claude-opus-4-7', object='chat.completion', system_fingerprint=None, choices=[Choices(finish_reason='stop', index=0, message=Message(content="Hello! I'm doing well, thanks for asking. I'm ready to help with whatever you needβ€”whether that's answering questions, brainstorming ideas, working through a problem, or just having a conversation. What can I do for you today?", role='assistant', tool_calls=None, function_call=None, provider_specific_fields={'citations': None, 'thinking_blocks': None}))], usage=Usage(completion_tokens=76, prompt_tokens=19, total_tokens=95, completion_tokens_details=CompletionTokensDetailsWrapper(accepted_prediction_tokens=None, audio_tokens=None, reasoning_tokens=0, rejected_prediction_tokens=None, text_tokens=76, image_tokens=None, video_tokens=None), prompt_tokens_details=PromptTokensDetailsWrapper(audio_tokens=None, cached_tokens=0, text_tokens=19, image_tokens=None, video_tokens=None, cache_creation_tokens=0, cache_creation_token_details=CacheCreationTokenDetails(ephemeral_5m_input_tokens=0, ephemeral_1h_input_tokens=0)), cache_creation_input_tokens=0, cache_read_input_tokens=0, inference_geo='global', speed=None)) -17:27:42 - LiteLLM:DEBUG: utils.py:482 - Async Wrapper: Completed Call, calling async_success_handler: > -17:27:42 - LiteLLM:DEBUG: litellm_logging.py:3312 - Filtered callbacks: [] - -#------------------------------------------------------------# -# # -# 'I get frustrated when the product...' # -# https://github.com/BerriAI/litellm/issues/new # -# # -#------------------------------------------------------------# - - Thank you for using LiteLLM! - Krrish & Ishaan - - - -Give Feedback / Get Help: https://github.com/BerriAI/litellm/issues/new - - -LiteLLM: Proxy initialized with Config, Set models: - gpt-realtime - claude-opus-4-7 -INFO: 127.0.0.1:51006 - "POST /v1/chat/completions HTTP/1.1" 200 OK -17:27:42 - LiteLLM:DEBUG: utils.py:482 - Logging Details LiteLLM-Async Success Call, cache_hit=None -17:27:42 - LiteLLM:DEBUG: utils.py:5620 - checking potential_model_names in litellm.model_cost: {'split_model': 'claude-opus-4-7', 'combined_model_name': 'anthropic/claude-opus-4-7', 'stripped_model_name': 'claude-opus-4-7', 'combined_stripped_model_name': 'anthropic/claude-opus-4-7', 'custom_llm_provider': 'anthropic'} -{ - "id": "chatcmpl-db0b1b06-c051-4281-98ff-6add4a98effb", - "trace_id": "066d410c-af0d-4e32-aca5-e3e657fbaff2", - "call_type": "acompletion", - "cache_hit": null, - "stream": null, - "status": "success", - "status_fields": { - "llm_api_status": "success", - "guardrail_status": "not_run" - }, - "custom_llm_provider": "anthropic", - "saved_cache_cost": 0.0, - "startTime": 1776686259.704125, - "endTime": 1776686262.173925, - "completionStartTime": 1776686262.173925, - "response_time": 2.4697999954223633, - "model": "anthropic/claude-opus-4-7", - "metadata": { - "user_api_key_hash": "88dc28d0f030c55ed4ab77ed8faf098196cb1c05df778539800c9f1243fe6b4b", - "user_api_key_alias": null, - "user_api_key_spend": 0.0, - "user_api_key_max_budget": null, - "user_api_key_budget_reset_at": null, - "user_api_key_team_id": null, - "user_api_key_org_id": null, - "user_api_key_org_alias": null, - "user_api_key_project_id": null, - "user_api_key_project_alias": null, - "user_api_key_user_id": "default_user_id", - "user_api_key_team_alias": null, - "user_api_key_user_email": null, - "user_api_key_end_user_id": null, - "user_api_key_request_route": "/v1/chat/completions", - "spend_logs_metadata": null, - "requester_ip_address": "127.0.0.1", - "user_agent": "PostmanRuntime/7.53.0", - "requester_metadata": { - "headers": { - "content-type": "application/json", - "user-agent": "PostmanRuntime/7.53.0", - "accept": "*/*", - "postman-token": "08d81a08-addf-4a27-a9f8-302f317d7296", - "host": "0.0.0.0:4000", - "accept-encoding": "gzip, deflate, br", - "connection": "keep-alive", - "content-length": "183" - } - }, - "prompt_management_metadata": null, - "applied_guardrails": [], - "mcp_tool_call_metadata": null, - "vector_store_request_metadata": null, - "usage_object": { - "completion_tokens": 76, - "prompt_tokens": 19, - "total_tokens": 95, - "completion_tokens_details": { - "reasoning_tokens": 0, - "text_tokens": 76 - }, - "prompt_tokens_details": { - "audio_tokens": null, - "cached_tokens": 0, - "text_tokens": 19, - "image_tokens": null, - "video_tokens": null, - "cache_creation_tokens": 0, - "cache_creation_token_details": { - "ephemeral_5m_input_tokens": 0, - "ephemeral_1h_input_tokens": 0 - } - }, - "cache_creation_input_tokens": 0, - "cache_read_input_tokens": 0, - "inference_geo": "global", - "speed": null - }, - "requester_custom_headers": {}, - "cold_storage_object_key": null, - "user_api_key_auth_metadata": {}, - "team_alias": null, - "team_id": null - }, - "cache_key": null, - "response_cost": 0.001995, - "cost_breakdown": { - "input_cost": 9.5e-05, - "output_cost": 0.0019, - "total_cost": 0.001995, - "tool_usage_cost": 0.0, - "original_cost": 0.001995, - "discount_percent": 0.0, - "discount_amount": 0.0, - "margin_percent": 0.0, - "margin_fixed_amount": 0.0, - "margin_total_amount": 0.0 - }, - "total_tokens": 95, - "prompt_tokens": 19, - "completion_tokens": 76, - "request_tags": [ - "User-Agent: PostmanRuntime", - "User-Agent: PostmanRuntime/7.53.0" - ], - "end_user": "", - "api_base": "https://api.anthropic.com/v1/messages", - "model_group": "claude-opus-4-7", - "model_id": "91a8591ffdee6edec7016031dc8a092b14ce59f85b1b105c0afa69280775d380", - "requester_ip_address": "127.0.0.1", - "user_agent": "PostmanRuntime/7.53.0", - "messages": [ - { - "role": "user", - "content": "Hello, how are you?" - } - ], - "response": { - "id": "chatcmpl-db0b1b06-c051-4281-98ff-6add4a98effb", - "created": 1776686262, - "model": "claude-opus-4-7", - "object": "chat.completion", - "system_fingerprint": null, - "choices": [ - { - "finish_reason": "stop", - "index": 0, - "message": { - "content": "Hello! I'm doing well, thanks for asking. I'm ready to help with whatever you need\u2014whether that's answering questions, brainstorming ideas, working through a problem, or just having a conversation. What can I do for you today?", - "role": "assistant", - "tool_calls": null, - "function_call": null, - "provider_specific_fields": { - "citations": null, - "thinking_blocks": null - } - } - } - ], - "usage": { - "completion_tokens": 76, - "prompt_tokens": 19, - "total_tokens": 95, - "completion_tokens_details": { - "reasoning_tokens": 0, - "text_tokens": 76 - }, - "prompt_tokens_details": { - "audio_tokens": null, - "cached_tokens": 0, - "text_tokens": 19, - "image_tokens": null, - "video_tokens": null, - "cache_creation_tokens": 0, - "cache_creation_token_details": { - "ephemeral_5m_input_tokens": 0, - "ephemeral_1h_input_tokens": 0 - } - }, - "cache_creation_input_tokens": 0, - "cache_read_input_tokens": 0, - "inference_geo": "global", - "speed": null - } - }, - "model_parameters": { - "stream": false, - "reasoning_effort": "max" - }, - "hidden_params": { - "model_id": "91a8591ffdee6edec7016031dc8a092b14ce59f85b1b105c0afa69280775d380", - "cache_key": null, - "api_base": null, - "response_cost": 0.001995, - "additional_headers": { - "x_ratelimit_limit_requests": 50, - "x_ratelimit_limit_tokens": 38000, - "x_ratelimit_remaining_requests": 49, - "x_ratelimit_remaining_tokens": 38000, - "llm_provider-date": "Mon, 20 Apr 2026 11:57:42 GMT", - "llm_provider-content-type": "application/json", - "llm_provider-transfer-encoding": "chunked", - "llm_provider-connection": "keep-alive", - "llm_provider-anthropic-ratelimit-input-tokens-limit": "30000", - "llm_provider-anthropic-ratelimit-input-tokens-remaining": "30000", - "llm_provider-anthropic-ratelimit-input-tokens-reset": "2026-04-20T11:57:41Z", - "llm_provider-anthropic-ratelimit-output-tokens-limit": "8000", - "llm_provider-anthropic-ratelimit-output-tokens-remaining": "8000", - "llm_provider-anthropic-ratelimit-output-tokens-reset": "2026-04-20T11:57:42Z", - "llm_provider-anthropic-ratelimit-requests-limit": "50", - "llm_provider-anthropic-ratelimit-requests-remaining": "49", - "llm_provider-anthropic-ratelimit-requests-reset": "2026-04-20T11:57:41Z", - "llm_provider-anthropic-ratelimit-tokens-limit": "38000", - "llm_provider-anthropic-ratelimit-tokens-remaining": "38000", - "llm_provider-anthropic-ratelimit-tokens-reset": "2026-04-20T11:57:41Z", - "llm_provider-request-id": "req_011CaEyx63z2yDaDF1967jyU", - "llm_provider-strict-transport-security": "max-age=31536000; includeSubDomains; preload", - "llm_provider-anthropic-organization-id": "cf4012a8-5914-4165-9573-d00b88e3d39e", - "llm_provider-server": "cloudflare", - "llm_provider-x-envoy-upstream-service-time": "2017", - "llm_provider-content-encoding": "gzip", - "llm_provider-vary": "Accept-Encoding", - "llm_provider-server-timing": "x-originResponse;dur=2019", - "llm_provider-cf-cache-status": "DYNAMIC", - "llm_provider-set-cookie": "_cfuvid=tNQ_aMy8wM6wZvYT8Gx.fxKqkk22mUq7vp0IzBG.r8g-1776686259.790975-1.0.1.1-eRorMbMzA43iO9c9lvJ8g5QwV1nQi.QPuE8FLWcVVeM; HttpOnly; SameSite=None; Secure; Path=/; Domain=api.anthropic.com", - "llm_provider-content-security-policy": "default-src 'none'; frame-ancestors 'none'", - "llm_provider-x-robots-tag": "none", - "llm_provider-cf-ray": "9ef3f903ab1347e7-BOM", - "llm_provider-x-ratelimit-limit-requests": "50", - "llm_provider-x-ratelimit-remaining-requests": "49", - "llm_provider-x-ratelimit-limit-tokens": "38000", - "llm_provider-x-ratelimit-remaining-tokens": "38000", - "x-litellm-model-group": "claude-opus-4-7", - "x-litellm-attempted-retries": 0, - "x-litellm-attempted-fallbacks": 0 - }, - "litellm_overhead_time_ms": 15.177, - "batch_models": null, - "litellm_model_name": "anthropic/claude-opus-4-7", - "usage_object": null - }, - "model_map_information": { - "model_map_key": "claude-opus-4-7", - "model_map_value": { - "key": "claude-opus-4-7", - "max_tokens": 128000, - "max_input_tokens": 1000000, - "max_output_tokens": 128000, - "input_cost_per_token": 5e-06, - "input_cost_per_token_flex": null, - "input_cost_per_token_priority": null, - "cache_creation_input_token_cost": 6.25e-06, - "cache_creation_input_token_cost_above_200k_tokens": null, - "cache_read_input_token_cost": 5e-07, - "cache_read_input_token_cost_above_200k_tokens": null, - "cache_read_input_token_cost_above_272k_tokens": null, - "cache_read_input_token_cost_flex": null, - "cache_read_input_token_cost_priority": null, - "cache_creation_input_token_cost_above_1hr": 1e-05, - "input_cost_per_character": null, - "input_cost_per_token_above_128k_tokens": null, - "input_cost_per_token_above_200k_tokens": null, - "input_cost_per_token_above_272k_tokens": null, - "input_cost_per_query": null, - "input_cost_per_second": null, - "input_cost_per_audio_token": null, - "input_cost_per_image_token": null, - "input_cost_per_image": null, - "input_cost_per_audio_per_second": null, - "input_cost_per_video_per_second": null, - "input_cost_per_token_batches": null, - "output_cost_per_token_batches": null, - "output_cost_per_token": 2.5e-05, - "output_cost_per_token_flex": null, - "output_cost_per_token_priority": null, - "output_cost_per_audio_token": null, - "output_cost_per_character": null, - "output_cost_per_reasoning_token": null, - "output_cost_per_token_above_128k_tokens": null, - "output_cost_per_character_above_128k_tokens": null, - "output_cost_per_token_above_200k_tokens": null, - "output_cost_per_token_above_272k_tokens": null, - "output_cost_per_second": null, - "output_cost_per_second_1080p": null, - "output_cost_per_video_per_second": null, - "output_cost_per_image": null, - "output_cost_per_image_token": null, - "output_vector_size": null, - "citation_cost_per_token": null, - "tiered_pricing": null, - "litellm_provider": "anthropic", - "mode": "chat", - "supports_system_messages": null, - "supports_response_schema": true, - "supports_vision": true, - "supports_function_calling": true, - "supports_tool_choice": true, - "supports_assistant_prefill": false, - "supports_prompt_caching": true, - "supports_audio_input": null, - "supports_audio_output": null, - "supports_pdf_input": true, - "supports_embedding_image_input": null, - "supports_native_streaming": null, - "supports_native_structured_output": null, - "supports_web_search": null, - "supports_url_context": null, - "supports_reasoning": true, - "supports_none_reasoning_effort": null, - "supports_xhigh_reasoning_effort": true, - "supports_computer_use": true, - "search_context_cost_per_query": { - "search_context_size_high": 0.01, - "search_context_size_low": 0.01, - "search_context_size_medium": 0.01 - }, - "tpm": null, - "rpm": null, - "ocr_cost_per_page": null, - "annotation_cost_per_page": null, - "provider_specific_entry": { - "us": 1.1, - "fast": 6.0 - }, - "uses_embed_content": null, - "supported_openai_params": [ - "stream", - "stop", - "temperature", - "top_p", - "max_tokens", - "max_completion_tokens", - "tools", - "tool_choice", - "extra_headers", - "parallel_tool_calls", - "response_format", - "user", - "web_search_options", - "speed", - "context_management", - "cache_control", - "thinking", - "reasoning_effort" - ] - } - }, - "error_str": null, - "error_information": { - "error_code": "", - "error_class": "", - "llm_provider": "", - "traceback": "", - "error_message": "" - }, - "response_cost_failure_debug_info": null, - "guardrail_information": null, - "standard_built_in_tools_params": { - "web_search_options": null, - "file_search": null - } -}17:27:42 - LiteLLM:DEBUG: callback_controls.py:37 - Dynamically disabled callbacks from x-litellm-disable-callbacks: None -17:27:42 - LiteLLM:DEBUG: callback_controls.py:38 - Checking if > is disabled via headers. Disable callbacks from headers: None -17:27:42 - LiteLLM:DEBUG: utils.py:5620 - checking potential_model_names in litellm.model_cost: {'split_model': 'claude-opus-4-7', 'combined_model_name': 'anthropic/claude-opus-4-7', 'stripped_model_name': 'anthropic/claude-opus-4-7', 'combined_stripped_model_name': 'anthropic/claude-opus-4-7', 'custom_llm_provider': 'anthropic'} -17:27:42 - LiteLLM:DEBUG: callback_controls.py:37 - Dynamically disabled callbacks from x-litellm-disable-callbacks: None -17:27:42 - LiteLLM:DEBUG: callback_controls.py:38 - Checking if is disabled via headers. Disable callbacks from headers: None -17:27:42 - LiteLLM Proxy:DEBUG: proxy_track_cost_callback.py:152 - INSIDE _PROXY_track_cost_callback -17:27:42 - LiteLLM Proxy:DEBUG: proxy_track_cost_callback.py:154 - kwargs stream: False + complete streaming response: None -17:27:42 - LiteLLM Proxy:DEBUG: proxy_track_cost_callback.py:186 - user_api_key 88dc28d0f030c55ed4ab77ed8faf098196cb1c05df778539800c9f1243fe6b4b, user_id default_user_id, team_id None, end_user_id None -17:27:42 - LiteLLM Proxy:DEBUG: db_spend_update_writer.py:119 - Enters prisma db call, response_cost: 0.001995, token: 88dc28d0f030c55ed4ab77ed8faf098196cb1c05df778539800c9f1243fe6b4b; user_id: default_user_id; team_id: None -17:27:42 - LiteLLM Proxy:DEBUG: spend_tracking_utils.py:116 - getting payload for SpendLogs, available keys in metadata: ['headers', 'requester_metadata', 'user_api_key_hash', 'user_api_key_alias', 'user_api_key_spend', 'user_api_key_max_budget', 'user_api_key_team_id', 'user_api_key_project_id', 'user_api_key_project_alias', 'user_api_key_user_id', 'user_api_key_org_id', 'user_api_key_org_alias', 'user_api_key_team_alias', 'user_api_key_end_user_id', 'user_api_key_user_email', 'user_api_key_request_route', 'user_api_key_budget_reset_at', 'user_api_key_auth_metadata', 'user_api_key', 'agent_id', 'user_api_end_user_max_budget', 'user_api_key_auth', 'litellm_api_version', 'global_max_parallel_requests', 'user_api_key_team_max_budget', 'user_api_key_team_spend', 'user_api_key_model_max_budget', 'user_api_key_end_user_model_max_budget', 'user_api_key_user_spend', 'user_api_key_user_max_budget', 'user_api_key_metadata', 'user_api_key_team_metadata', 'user_api_key_object_permission_id', 'user_api_key_team_object_permission_id', 'endpoint', 'litellm_parent_otel_span', 'requester_ip_address', 'user_agent', 'queue_time_seconds', 'model_group', 'model_group_alias', 'model_group_size', 'attempted_retries', 'max_retries', 'deployment', 'model_info', 'api_base', 'deployment_model_name', 'caching_groups', 'hidden_params'] -17:27:42 - LiteLLM Proxy:DEBUG: spend_tracking_utils.py:480 - SpendTable: created payload - request_id: chatcmpl-db0b1b06-c051-4281-98ff-6add4a98effb, model: anthropic/claude-opus-4-7, spend: 0.001995 -17:27:42 - LiteLLM Proxy:DEBUG: db_spend_update_writer.py:726 - Writing spend log to db - request_id: chatcmpl-db0b1b06-c051-4281-98ff-6add4a98effb, spend: 0.001995 -17:27:42 - LiteLLM Proxy:DEBUG: db_spend_update_writer.py:193 - Runs spend update on all tables -17:27:42 - LiteLLM:DEBUG: callback_controls.py:37 - Dynamically disabled callbacks from x-litellm-disable-callbacks: None -17:27:42 - LiteLLM:DEBUG: callback_controls.py:38 - Checking if is disabled via headers. Disable callbacks from headers: None -17:27:42 - LiteLLM:DEBUG: callback_controls.py:37 - Dynamically disabled callbacks from x-litellm-disable-callbacks: None -17:27:42 - LiteLLM:DEBUG: callback_controls.py:38 - Checking if is disabled via headers. Disable callbacks from headers: None -17:27:42 - LiteLLM Proxy:DEBUG: model_max_budget_limiter.py:227 - in RouterBudgetLimiting.async_log_success_event -17:27:42 - LiteLLM Proxy:DEBUG: model_max_budget_limiter.py:252 - Not running _PROXY_VirtualKeyModelMaxBudgetLimiter.async_log_success_event because user_api_key_model_max_budget and user_api_key_end_user_model_max_budget are None or empty. -17:27:42 - LiteLLM:DEBUG: callback_controls.py:37 - Dynamically disabled callbacks from x-litellm-disable-callbacks: None -17:27:42 - LiteLLM:DEBUG: callback_controls.py:38 - Checking if is disabled via headers. Disable callbacks from headers: None -17:27:42 - LiteLLM:DEBUG: callback_controls.py:37 - Dynamically disabled callbacks from x-litellm-disable-callbacks: None -17:27:42 - LiteLLM:DEBUG: callback_controls.py:38 - Checking if is disabled via headers. Disable callbacks from headers: None -17:27:42 - LiteLLM Proxy:DEBUG: parallel_request_limiter_v3.py:1740 - INSIDE parallel request limiter ASYNC SUCCESS LOGGING -17:27:42 - LiteLLM Proxy:DEBUG: parallel_request_limiter_v3.py:1506 - TTL preservation script not available, using regular pipeline -17:27:42 - LiteLLM:DEBUG: callback_controls.py:37 - Dynamically disabled callbacks from x-litellm-disable-callbacks: None -17:27:42 - LiteLLM:DEBUG: callback_controls.py:38 - Checking if is disabled via headers. Disable callbacks from headers: None -17:27:42 - LiteLLM:DEBUG: callback_controls.py:37 - Dynamically disabled callbacks from x-litellm-disable-callbacks: None -17:27:42 - LiteLLM:DEBUG: callback_controls.py:38 - Checking if is disabled via headers. Disable callbacks from headers: None -17:27:42 - LiteLLM:DEBUG: callback_controls.py:37 - Dynamically disabled callbacks from x-litellm-disable-callbacks: None -17:27:42 - LiteLLM:DEBUG: callback_controls.py:38 - Checking if is disabled via headers. Disable callbacks from headers: None -17:27:42 - LiteLLM:DEBUG: callback_controls.py:37 - Dynamically disabled callbacks from x-litellm-disable-callbacks: None -17:27:42 - LiteLLM:DEBUG: callback_controls.py:38 - Checking if is disabled via headers. Disable callbacks from headers: None -17:27:42 - LiteLLM:DEBUG: callback_controls.py:37 - Dynamically disabled callbacks from x-litellm-disable-callbacks: None -17:27:42 - LiteLLM:DEBUG: callback_controls.py:38 - Checking if is disabled via headers. Disable callbacks from headers: None -17:27:42 - LiteLLM:DEBUG: callback_controls.py:37 - Dynamically disabled callbacks from x-litellm-disable-callbacks: None -17:27:42 - LiteLLM:DEBUG: callback_controls.py:38 - Checking if is disabled via headers. Disable callbacks from headers: None -17:27:42 - LiteLLM:DEBUG: callback_controls.py:37 - Dynamically disabled callbacks from x-litellm-disable-callbacks: None -17:27:42 - LiteLLM:DEBUG: callback_controls.py:38 - Checking if is disabled via headers. Disable callbacks from headers: None -17:27:42 - LiteLLM Proxy:DEBUG: spend_update_queue.py:44 - Adding update to queue: {'entity_type': , 'entity_id': 'default_user_id', 'response_cost': 0.001995} -17:27:42 - LiteLLM Proxy:DEBUG: spend_update_queue.py:44 - Adding update to queue: {'entity_type': , 'entity_id': '88dc28d0f030c55ed4ab77ed8faf098196cb1c05df778539800c9f1243fe6b4b', 'response_cost': 0.001995} -17:27:42 - LiteLLM Proxy:DEBUG: db_spend_update_writer.py:562 - track_cost_callback: team_id is None or prisma_client is None. Not tracking spend for team -17:27:42 - LiteLLM Proxy:DEBUG: db_spend_update_writer.py:616 - track_cost_callback: org_id is None or prisma_client is None. Not tracking spend for org -17:27:42 - LiteLLM Proxy:DEBUG: spend_update_queue.py:44 - Adding update to queue: {'entity_type': , 'entity_id': 'User-Agent: PostmanRuntime', 'response_cost': 0.001995} -17:27:42 - LiteLLM Proxy:DEBUG: spend_update_queue.py:44 - Adding update to queue: {'entity_type': , 'entity_id': 'User-Agent: PostmanRuntime/7.53.0', 'response_cost': 0.001995} -17:27:42 - LiteLLM Proxy:DEBUG: db_spend_update_writer.py:1950 - Logged request status: success -17:27:42 - LiteLLM Proxy:DEBUG: daily_spend_update_queue.py:63 - Adding update to queue: {'default_user_id_2026-04-20_88dc28d0f030c55ed4ab77ed8faf098196cb1c05df778539800c9f1243fe6b4b_anthropic/claude-opus-4-7_anthropic_/chat/completions': {'user_id': 'default_user_id', 'date': '2026-04-20', 'REDACTED', 'model': 'anthropic/claude-opus-4-7', 'model_group': 'claude-opus-4-7', 'mcp_namespaced_tool_name': None, 'custom_llm_provider': 'anthropic', 'endpoint': '/chat/completions', 'prompt_tokens': 19, 'completion_tokens': 76, 'spend': 0.001995, 'api_requests': 1, 'successful_requests': 1, 'failed_requests': 0, 'cache_read_input_tokens': 0, 'cache_creation_input_tokens': 0}} -17:27:42 - LiteLLM Proxy:DEBUG: db_spend_update_writer.py:2119 - end_user is None or empty for request. Skipping incrementing end user spend. -17:27:42 - LiteLLM Proxy:DEBUG: db_spend_update_writer.py:1950 - Logged request status: success -17:27:42 - LiteLLM Proxy:DEBUG: daily_spend_update_queue.py:63 - Adding update to queue: {'_2026-04-20_88dc28d0f030c55ed4ab77ed8faf098196cb1c05df778539800c9f1243fe6b4b_anthropic/claude-opus-4-7_anthropic_/chat/completions': {'team_id': '', 'date': '2026-04-20', 'REDACTED', 'model': 'anthropic/claude-opus-4-7', 'model_group': 'claude-opus-4-7', 'mcp_namespaced_tool_name': None, 'custom_llm_provider': 'anthropic', 'endpoint': '/chat/completions', 'prompt_tokens': 19, 'completion_tokens': 76, 'spend': 0.001995, 'api_requests': 1, 'successful_requests': 1, 'failed_requests': 0, 'cache_read_input_tokens': 0, 'cache_creation_input_tokens': 0}} -17:27:42 - LiteLLM Proxy:DEBUG: db_spend_update_writer.py:2076 - organization_id is None for request. Skipping incrementing organization spend. -17:27:42 - LiteLLM Proxy:DEBUG: db_spend_update_writer.py:1950 - Logged request status: success -17:27:42 - LiteLLM Proxy:DEBUG: daily_spend_update_queue.py:63 - Adding update to queue: {'User-Agent: PostmanRuntime_2026-04-20_88dc28d0f030c55ed4ab77ed8faf098196cb1c05df778539800c9f1243fe6b4b_anthropic/claude-opus-4-7_anthropic_/chat/completions': {'tag': 'User-Agent: PostmanRuntime', 'date': '2026-04-20', 'REDACTED', 'model': 'anthropic/claude-opus-4-7', 'model_group': 'claude-opus-4-7', 'mcp_namespaced_tool_name': None, 'custom_llm_provider': 'anthropic', 'endpoint': '/chat/completions', 'prompt_tokens': 19, 'completion_tokens': 76, 'spend': 0.001995, 'api_requests': 1, 'successful_requests': 1, 'failed_requests': 0, 'cache_read_input_tokens': 0, 'cache_creation_input_tokens': 0, 'request_id': 'chatcmpl-db0b1b06-c051-4281-98ff-6add4a98effb'}} -17:27:42 - LiteLLM Proxy:DEBUG: daily_spend_update_queue.py:63 - Adding update to queue: {'User-Agent: PostmanRuntime/7.53.0_2026-04-20_88dc28d0f030c55ed4ab77ed8faf098196cb1c05df778539800c9f1243fe6b4b_anthropic/claude-opus-4-7_anthropic_/chat/completions': {'tag': 'User-Agent: PostmanRuntime/7.53.0', 'date': '2026-04-20', 'REDACTED', 'model': 'anthropic/claude-opus-4-7', 'model_group': 'claude-opus-4-7', 'mcp_namespaced_tool_name': None, 'custom_llm_provider': 'anthropic', 'endpoint': '/chat/completions', 'prompt_tokens': 19, 'completion_tokens': 76, 'spend': 0.001995, 'api_requests': 1, 'successful_requests': 1, 'failed_requests': 0, 'cache_read_input_tokens': 0, 'cache_creation_input_tokens': 0, 'request_id': 'chatcmpl-db0b1b06-c051-4281-98ff-6add4a98effb'}} -17:27:42 - LiteLLM Proxy:DEBUG: proxy_server.py:1945 - _update_key_cache: hashed_token=88dc28d0f030c55ed4ab77ed8faf098196cb1c05df778539800c9f1243fe6b4b -17:27:42 - LiteLLM Proxy:DEBUG: proxy_server.py:1947 - _update_key_cache: existing_spend_obj=token='88dc28d0f030c55ed4ab77ed8faf098196cb1c05df778539800c9f1243fe6b4b' key_name=None key_alias=None spend=0.0 max_budget=None expires=None models=[] aliases={} config={} user_id='default_user_id' team_id=None agent_id=None project_id=None max_parallel_requests=None metadata={} tpm_limit=None rpm_limit=None budget_duration=None budget_reset_at=None allowed_cache_controls=[] allowed_routes=[] permissions={} model_spend={} model_max_budget={} soft_budget_cooldown=False blocked=None litellm_budget_table=None org_id=None created_at=None created_by=None updated_at=None updated_by=None last_active=None object_permission_id=None object_permission=None access_group_ids=None rotation_count=0 auto_rotate=False rotation_interval=None last_rotation_at=None key_rotation_at=None router_settings=None budget_limits=None team_spend=None team_alias=None team_tpm_limit=None team_rpm_limit=None team_max_budget=None team_soft_budget=None team_models=[] team_blocked=False soft_budget=None team_model_aliases=None team_member=None team_metadata=None team_object_permission_id=None team_member_spend=None team_member_tpm_limit=None team_member_rpm_limit=None end_user_id=None end_user_tpm_limit=None end_user_rpm_limit=None end_user_max_budget=None end_user_model_max_budget=None organization_alias=None organization_max_budget=None organization_tpm_limit=None organization_rpm_limit=None organization_metadata=None project_alias=None project_metadata=None last_refreshed_at=1776686259.698418 REDACTED' user_role= allowed_model_region=None parent_otel_span=None rpm_limit_per_model=None tpm_limit_per_model=None user_tpm_limit=None user_rpm_limit=None user_email=None user_spend=None user_max_budget=None request_route='/v1/chat/completions' user=None created_by_user=None end_user_object_permission=None jwt_claims=None -17:27:47 - LiteLLM Proxy:DEBUG: utils.py:5096 - Spend logs queue size (1) below threshold (100), processing with backoff -17:27:47 - LiteLLM Proxy:INFO: utils.py:4813 - Spend tracking - processing 1 spend logs for DB write -17:27:47 - LiteLLM Proxy:DEBUG: utils.py:4850 - Flushed 1 logs to the DB. -17:27:47 - LiteLLM Proxy:DEBUG: utils.py:4859 - 1 logs processed. Remaining in queue: 0 -17:27:47 - LiteLLM Proxy:INFO: spend_update_queue.py:35 - Spend tracking - flushed 4 spend update items from in-memory queue -17:27:47 - LiteLLM Proxy:DEBUG: spend_update_queue.py:39 - Aggregating updates by entity type: [{'entity_type': , 'entity_id': 'default_user_id', 'response_cost': 0.001995}, {'entity_type': , 'entity_id': '88dc28d0f030c55ed4ab77ed8faf098196cb1c05df778539800c9f1243fe6b4b', 'response_cost': 0.001995}, {'entity_type': , 'entity_id': 'User-Agent: PostmanRuntime', 'response_cost': 0.001995}, {'entity_type': , 'entity_id': 'User-Agent: PostmanRuntime/7.53.0', 'response_cost': 0.001995}] -17:27:47 - LiteLLM Proxy:DEBUG: db_spend_update_writer.py:1122 - User Spend transactions: {'default_user_id': 0.001995} -17:27:47 - LiteLLM Proxy:DEBUG: db_spend_update_writer.py:1165 - End-User Spend transactions: {} -17:27:47 - LiteLLM Proxy:DEBUG: db_spend_update_writer.py:1180 - KEY Spend transactions: {'88dc28d0f030c55ed4ab77ed8faf098196cb1c05df778539800c9f1243fe6b4b': 0.001995} -17:27:47 - LiteLLM Proxy:DEBUG: db_spend_update_writer.py:1221 - Team Spend transactions: {} -17:27:47 - LiteLLM Proxy:DEBUG: db_spend_update_writer.py:1269 - Team Membership Spend transactions: {} -17:27:47 - LiteLLM Proxy:DEBUG: db_spend_update_writer.py:1339 - Org Spend transactions: {} -17:27:47 - LiteLLM Proxy:DEBUG: db_spend_update_writer.py:1432 - Tag Spend transactions: {'User-Agent: PostmanRuntime': 0.001995, 'User-Agent: PostmanRuntime/7.53.0': 0.001995} -17:27:47 - LiteLLM Proxy:DEBUG: db_spend_update_writer.py:1442 - Updating spend for Tag tag_name=User-Agent: PostmanRuntime by 0.001995 -17:27:47 - LiteLLM Proxy:DEBUG: db_spend_update_writer.py:1442 - Updating spend for Tag tag_name=User-Agent: PostmanRuntime/7.53.0 by 0.001995 -17:27:47 - LiteLLM Proxy:DEBUG: db_spend_update_writer.py:1432 - Agent Spend transactions: {} -17:27:47 - LiteLLM Proxy:INFO: daily_spend_update_queue.py:90 - Spend tracking - flushed 1 daily spend update items from in-memory queue -17:27:47 - LiteLLM Proxy:DEBUG: daily_spend_update_queue.py:99 - Aggregated daily spend update transactions: {'default_user_id_2026-04-20_88dc28d0f030c55ed4ab77ed8faf098196cb1c05df778539800c9f1243fe6b4b_anthropic/claude-opus-4-7_anthropic_/chat/completions': {'user_id': 'default_user_id', 'date': '2026-04-20', 'REDACTED', 'model': 'anthropic/claude-opus-4-7', 'model_group': 'claude-opus-4-7', 'mcp_namespaced_tool_name': None, 'custom_llm_provider': 'anthropic', 'endpoint': '/chat/completions', 'prompt_tokens': 19, 'completion_tokens': 76, 'spend': 0.001995, 'api_requests': 1, 'successful_requests': 1, 'failed_requests': 0, 'cache_read_input_tokens': 0, 'cache_creation_input_tokens': 0}} -17:27:47 - LiteLLM Proxy:DEBUG: db_spend_update_writer.py:1573 - Daily User Spend transactions: 1 -17:27:47 - LiteLLM Proxy:DEBUG: db_spend_update_writer.py:1744 - Processed 1 daily user transactions in 0.02s -17:27:47 - LiteLLM Proxy:INFO: daily_spend_update_queue.py:90 - Spend tracking - flushed 1 daily spend update items from in-memory queue -17:27:47 - LiteLLM Proxy:DEBUG: daily_spend_update_queue.py:99 - Aggregated daily spend update transactions: {'_2026-04-20_88dc28d0f030c55ed4ab77ed8faf098196cb1c05df778539800c9f1243fe6b4b_anthropic/claude-opus-4-7_anthropic_/chat/completions': {'team_id': '', 'date': '2026-04-20', 'REDACTED', 'model': 'anthropic/claude-opus-4-7', 'model_group': 'claude-opus-4-7', 'mcp_namespaced_tool_name': None, 'custom_llm_provider': 'anthropic', 'endpoint': '/chat/completions', 'prompt_tokens': 19, 'completion_tokens': 76, 'spend': 0.001995, 'api_requests': 1, 'successful_requests': 1, 'failed_requests': 0, 'cache_read_input_tokens': 0, 'cache_creation_input_tokens': 0}} -17:27:47 - LiteLLM Proxy:DEBUG: db_spend_update_writer.py:1573 - Daily Team Spend transactions: 1 -17:27:47 - LiteLLM Proxy:DEBUG: db_spend_update_writer.py:1744 - Processed 1 daily team transactions in 0.01s -17:27:47 - LiteLLM Proxy:DEBUG: daily_spend_update_queue.py:99 - Aggregated daily spend update transactions: {} -17:27:47 - LiteLLM Proxy:DEBUG: db_spend_update_writer.py:1573 - Daily Org Spend transactions: 0 -17:27:47 - LiteLLM Proxy:DEBUG: db_spend_update_writer.py:1606 - No new transactions to process for daily org spend update -17:27:47 - LiteLLM Proxy:DEBUG: daily_spend_update_queue.py:99 - Aggregated daily spend update transactions: {} -17:27:47 - LiteLLM Proxy:DEBUG: db_spend_update_writer.py:1573 - Daily End_user Spend transactions: 0 -17:27:47 - LiteLLM Proxy:DEBUG: db_spend_update_writer.py:1606 - No new transactions to process for daily end_user spend update -17:27:47 - LiteLLM Proxy:DEBUG: daily_spend_update_queue.py:99 - Aggregated daily spend update transactions: {} -17:27:47 - LiteLLM Proxy:DEBUG: db_spend_update_writer.py:1573 - Daily Agent Spend transactions: 0 -17:27:47 - LiteLLM Proxy:DEBUG: db_spend_update_writer.py:1606 - No new transactions to process for daily agent spend update -17:27:47 - LiteLLM Proxy:DEBUG: utils.py:4931 - Spend Logs transactions: 0 -17:28:01 - LiteLLM Proxy:ERROR: utils.py:3918 - prisma-query-engine PID 66228 exited (waitpid thread); triggering reconnect. -17:28:01 - LiteLLM Proxy:WARNING: utils.py:4178 - Attempting Prisma DB reconnect. reason=engine_process_death -17:28:01 - LiteLLM Proxy:WARNING: utils.py:4107 - prisma-query-engine PID 0 is dead; reconnecting. -query-engine ac9d7041ed77bcc8a8dbd2ab6616b39013829574 -INFO: Shutting down -INFO: Waiting for application shutdown. -17:28:01 - LiteLLM Proxy:INFO: proxy_server.py:971 - SESSION REUSE: Closed shared aiohttp session -17:28:01 - LiteLLM Proxy:DEBUG: utils.py:4082 - Stopped engine process watcher. -17:28:01 - LiteLLM Proxy:INFO: utils.py:4322 - Stopped Prisma DB health watchdog -17:28:01 - LiteLLM Proxy:INFO: proxy_server.py:708 - Shutting down LiteLLM Proxy Server -17:28:01 - LiteLLM Proxy:DEBUG: proxy_server.py:710 - Disconnecting from Prisma -INFO: Application shutdown complete. -INFO: Finished server process [66039] -17:28:01 - LiteLLM:DEBUG: logging_worker.py:134 - LoggingWorker cancelled during shutdown -17:28:01 - LiteLLM:DEBUG: logging_worker.py:446 - [LoggingWorker] atexit: Queue is empty -