Add greenpt/glm-5.3-flash (/usr/bin/bash.11//usr/bin/bash.44 per 1M)

This commit is contained in:
inference-bot committed 2026-09-24 10:43:14 -06:00
1 parent f86e005062
commit 2369c3e848
1 file changed
+12
+12
View File
@@ -73,6 +73,18 @@ model_list:
input_cost_per_token: 0.00000025
output_cost_per_token: 0.00000080
# GLM-5.3 Flash on GreenPT (note the dotted slug — distinct from Tinfoil's
# glm-5-3-flash). $0.11 in / $0.44 out / $0.022 cached per 1M.
- model_name: greenpt/glm-5.3-flash
litellm_params:
model: openai/glm-5.3-flash
api_base: https://api.greenpt.ai/v1
api_key: os.environ/GREENPT_API_KEY
model_info:
input_cost_per_token: 0.00000011
output_cost_per_token: 0.00000044
cache_read_input_token_cost: 0.000000022
# Fallback: if DeepSeek is down, use GPT-OSS
# NOTE: DeepSeek V4.1 Flash is a reasoning model — generations routinely run
# longer than 30s. A 30s timeout caused LiteLLM to abandon in-flight DeepSeek