test: move the remaining live groq call sites off the retired llama models (#37426)

The earlier sweep only caught the conformance suite in tests/llm_translation.
Groq retired llama-3.1-8b-instant alongside llama-3.3-70b-versatile, and four
tests under tests/local_testing still call them for real, so litellm_router_testing
and both local_testing shards 404 with model_not_found.

Only the sites that leave the process move. The chunk fixtures in
test_stream_chunk_builder, and the cost and routing tests that never open a
socket, keep the old ids because the string is data there, not a request.
This commit is contained in:
yuneng-jiang 2026-08-18 20:26:17 -07:00 • committed by GitHub
parent 16bc32fa23
commit bcead282e2
No known key found for this signature in database
GPG key ID: B5690EEEBB952194
4 changed files with 7 additions and 7 deletions

View file

@ -15,7 +15,7 @@ from litellm import completion, embedding
litellm.set_verbose = True
model_alias_map = {"good-model": "groq/llama-3.1-8b-instant"}
model_alias_map = {"good-model": "groq/openai/gpt-oss-120b"}
def test_model_alias_map(caplog):
@ -34,7 +34,7 @@ def test_model_alias_map(caplog):
if rec.levelname == "ERROR" and rec.name.startswith("LiteLLM"):
pytest.fail(f"Unexpected litellm ERROR log: {rec.getMessage()}")
assert "llama-3.1-8b-instant" in response.model
assert "gpt-oss-120b" in response.model
except litellm.ServiceUnavailableError:
pass
except Exception as e:

View file

@ -120,7 +120,7 @@ async def test_router_provider_wildcard_routing():
print("response 2 = ", response2)
response3 = await router.acompletion(
model="groq/llama-3.1-8b-instant",
model="groq/openai/gpt-oss-120b",
messages=[{"role": "user", "content": "hello"}],
)

View file

@ -44,7 +44,7 @@ async def test_batch_completion_multiple_models(mode):
{
"model_name": "groq-llama",
"litellm_params": {
"model": "groq/llama-3.1-8b-instant",
"model": "groq/openai/gpt-oss-120b",
},
},
]
@ -143,7 +143,7 @@ async def test_batch_completion_fastest_response_streaming():
{
"model_name": "groq-llama",
"litellm_params": {
"model": "groq/llama-3.1-8b-instant",
"model": "groq/openai/gpt-oss-120b",
},
},
]
@ -179,7 +179,7 @@ async def test_batch_completion_multiple_models_multiple_messages():
{
"model_name": "groq-llama",
"litellm_params": {
"model": "groq/llama-3.1-8b-instant",
"model": "groq/openai/gpt-oss-120b",
},
},
]

View file

@ -871,7 +871,7 @@ def load_env():
}
LLAMA3_3 = {
"messages": messages,
"model": "groq/llama-3.3-70b-versatile",
"model": "groq/openai/gpt-oss-120b",
"api_base": "https://api.groq.com/openai/v1",
"temperature": 0.0,
"tools": tools,