mirror of
https://github.com/RooVetGit/Roo-Code.git
synced 2026-10-06 02:47:56 +00:00
feat: add zai-glm-4.6 model to Cerebras and set gpt-oss-120b as default (#8920)
* feat: add zai-glm-4.6 model and update gpt-oss-120b for Cerebras - Add zai-glm-4.6 with 128K context window and 40K max tokens - Set zai-glm-4.6 as default Cerebras model - Update gpt-oss-120b to 128K context and 40K max tokens * feat: add zai-glm-4.6 model to Cerebras provider - Add zai-glm-4.6 with 128K context window and 40K max tokens - Set zai-glm-4.6 as default Cerebras model - Model provides ~2000 tokens/s for general-purpose tasks * add [SOON TO BE DEPRECATED] warning for Q3C * chore: set gpt-oss-120b as default Cerebras model * Fix cerebras test: update expected default model to gpt-oss-120b * Apply suggestion from @mrubens Co-authored-by: Matt Rubens <mrubens@users.noreply.github.com> --------- Co-authored-by: kevint-cerebras <kevin.taylor@cerebras.net> Co-authored-by: Matt Rubens <mrubens@users.noreply.github.com>
This commit is contained in:
parent
bd5807bac1
commit
ed45d1c081
2 changed files with 13 additions and 4 deletions
|
|
@ -3,9 +3,18 @@ import type { ModelInfo } from "../model.js"
|
|||
// https://inference-docs.cerebras.ai/api-reference/chat-completions
|
||||
export type CerebrasModelId = keyof typeof cerebrasModels
|
||||
|
||||
export const cerebrasDefaultModelId: CerebrasModelId = "qwen-3-coder-480b-free"
|
||||
export const cerebrasDefaultModelId: CerebrasModelId = "gpt-oss-120b"
|
||||
|
||||
export const cerebrasModels = {
|
||||
"zai-glm-4.6": {
|
||||
maxTokens: 16_384,
|
||||
contextWindow: 128000,
|
||||
supportsImages: false,
|
||||
supportsPromptCache: false,
|
||||
inputPrice: 0,
|
||||
outputPrice: 0,
|
||||
description: "Highly intelligent general-purpose model with ~2000 tokens/s",
|
||||
},
|
||||
"qwen-3-coder-480b-free": {
|
||||
maxTokens: 40000,
|
||||
contextWindow: 64000,
|
||||
|
|
@ -14,7 +23,7 @@ export const cerebrasModels = {
|
|||
inputPrice: 0,
|
||||
outputPrice: 0,
|
||||
description:
|
||||
"SOTA coding model with ~2000 tokens/s ($0 free tier)\n\n• Use this if you don't have a Cerebras subscription\n• 64K context window\n• Rate limits: 150K TPM, 1M TPH/TPD, 10 RPM, 100 RPH/RPD\n\nUpgrade for higher limits: [https://cloud.cerebras.ai/?utm=roocode](https://cloud.cerebras.ai/?utm=roocode)",
|
||||
"[SOON TO BE DEPRECATED] SOTA coding model with ~2000 tokens/s ($0 free tier)\n\n• Use this if you don't have a Cerebras subscription\n• 64K context window\n• Rate limits: 150K TPM, 1M TPH/TPD, 10 RPM, 100 RPH/RPD\n\nUpgrade for higher limits: [https://cloud.cerebras.ai/?utm=roocode](https://cloud.cerebras.ai/?utm=roocode)",
|
||||
},
|
||||
"qwen-3-coder-480b": {
|
||||
maxTokens: 40000,
|
||||
|
|
@ -24,7 +33,7 @@ export const cerebrasModels = {
|
|||
inputPrice: 0,
|
||||
outputPrice: 0,
|
||||
description:
|
||||
"SOTA coding model with ~2000 tokens/s ($50/$250 paid tiers)\n\n• Use this if you have a Cerebras subscription\n• 131K context window with higher rate limits",
|
||||
"[SOON TO BE DEPRECATED] SOTA coding model with ~2000 tokens/s ($50/$250 paid tiers)\n\n• Use this if you have a Cerebras subscription\n• 131K context window with higher rate limits",
|
||||
},
|
||||
"qwen-3-235b-a22b-instruct-2507": {
|
||||
maxTokens: 64000,
|
||||
|
|
|
|||
|
|
@ -56,7 +56,7 @@ describe("CerebrasHandler", () => {
|
|||
it("should fallback to default model when apiModelId is not provided", () => {
|
||||
const handlerWithoutModel = new CerebrasHandler({ cerebrasApiKey: "test" })
|
||||
const { id } = handlerWithoutModel.getModel()
|
||||
expect(id).toBe("qwen-3-coder-480b") // cerebrasDefaultModelId (routed)
|
||||
expect(id).toBe("gpt-oss-120b") // cerebrasDefaultModelId
|
||||
})
|
||||
})
|
||||
|
||||
|
|
|
|||
Loading…
Add table
Reference in a new issue