Add greenpt/glm-5.3-flash (/usr/bin/bash.11//usr/bin/bash.44 per 1M)
This commit is contained in:
1 parent
f86e005062
commit
2369c3e848
1 file changed
+12
+12
@@ -73,6 +73,18 @@ model_list:
|
|||||||
input_cost_per_token: 0.00000025
|
input_cost_per_token: 0.00000025
|
||||||
output_cost_per_token: 0.00000080
|
output_cost_per_token: 0.00000080
|
||||||
|
|
||||||
|
# GLM-5.3 Flash on GreenPT (note the dotted slug — distinct from Tinfoil's
|
||||||
|
# glm-5-3-flash). $0.11 in / $0.44 out / $0.022 cached per 1M.
|
||||||
|
- model_name: greenpt/glm-5.3-flash
|
||||||
|
litellm_params:
|
||||||
|
model: openai/glm-5.3-flash
|
||||||
|
api_base: https://api.greenpt.ai/v1
|
||||||
|
api_key: os.environ/GREENPT_API_KEY
|
||||||
|
model_info:
|
||||||
|
input_cost_per_token: 0.00000011
|
||||||
|
output_cost_per_token: 0.00000044
|
||||||
|
cache_read_input_token_cost: 0.000000022
|
||||||
|
|
||||||
# Fallback: if DeepSeek is down, use GPT-OSS
|
# Fallback: if DeepSeek is down, use GPT-OSS
|
||||||
# NOTE: DeepSeek V4.1 Flash is a reasoning model — generations routinely run
|
# NOTE: DeepSeek V4.1 Flash is a reasoning model — generations routinely run
|
||||||
# longer than 30s. A 30s timeout caused LiteLLM to abandon in-flight DeepSeek
|
# longer than 30s. A 30s timeout caused LiteLLM to abandon in-flight DeepSeek
|
||||||
|
|||||||
Reference in new issue
Block a user