Merge pull request #638 from deepinfra/deepinfra-models

deepinfra: Add supported models
This commit is contained in:
Ishaan Jaff 2023-10-20 09:10:32 -07:00 • committed by GitHub
commit 9d1ce893b7
No known key found for this signature in database
GPG key ID: 4AEE18F83AFDEB23
2 changed files with 8 additions and 3 deletions

View file

@ -99,6 +99,7 @@ ai21_models: List = []
nlp_cloud_models: List = []
aleph_alpha_models: List = []
bedrock_models: List = []
deepinfra_models: List = []
for key, value in model_cost.items():
if value.get('litellm_provider') == 'openai':
open_ai_chat_completion_models.append(key)
@ -127,6 +128,8 @@ for key, value in model_cost.items():
aleph_alpha_models.append(key)
elif value.get('litellm_provider') == 'bedrock':
bedrock_models.append(key)
elif value.get('litellm_provider') == 'deepinfra':
deepinfra_models.append(key)
# known openai compatible endpoints - we'll eventually move this list to the model_prices_and_context_window.json dictionary
openai_compatible_endpoints: List = [
@ -230,6 +233,7 @@ model_list = (
+ nlp_cloud_models
+ ollama_models
+ bedrock_models
+ deepinfra_models
)
provider_list: List = [
@ -271,6 +275,7 @@ models_by_provider: dict = {
"bedrock": bedrock_models,
"petals": petals_models,
"ollama": ollama_models,
"deepinfra": deepinfra_models,
}
# mapping for those models which have larger equivalents

View file

@ -572,9 +572,9 @@
"mode": "completion"
},
"deepinfra/meta-llama/Llama-2-70b-chat-hf": {
"max_tokens": 6144,
"input_cost_per_token": 0.000001875,
"output_cost_per_token": 0.000001875,
"max_tokens": 4096,
"input_cost_per_token": 0.000000700,
"output_cost_per_token": 0.000000950,
"litellm_provider": "deepinfra",
"mode": "chat"
},