Skip to content
Merged
Show file tree
Hide file tree
Changes from all commits
Commits
File filter

Filter by extension

Filter by extension

Conversations
Failed to load comments.
Loading
Jump to
Jump to file
Failed to load files.
Loading
Diff view
Diff view
16 changes: 16 additions & 0 deletions providers/greenpt/models/deepseek-v4-flash-0731.toml
Original file line number Diff line number Diff line change
@@ -0,0 +1,16 @@
# Cost: converted from GreenPT's EUR list price at 1.14 USD/EUR (rate captured 2026-08-01).
# Sources: https://docs.greenpt.ai/pricing and https://docs.greenpt.ai/model-cards
# reasoning_effort: thinking is enabled by default and "none" disables it.
# Values verified against the production API. See https://docs.greenpt.ai/chat-completion
# cache_read: discounted rate for cached prompt tokens. Cache writes are not
# charged. See https://docs.greenpt.ai/prompt-caching
base_model = "deepseek/deepseek-v4-flash"
name = "DeepSeek V4 Flash 0731"
release_date = "2026-07-31"
last_updated = "2026-07-31"
reasoning_options = [{ type = "effort", values = ["none", "minimal", "low", "medium", "high"] }]

[cost]
input = 0.1596
cache_read = 0.0456
output = 0.399
4 changes: 2 additions & 2 deletions providers/greenpt/models/gemma4.toml
Original file line number Diff line number Diff line change
@@ -1,7 +1,7 @@
# Cost: converted from GreenPT's EUR list price at 1.14 USD/EUR (rate captured 2026-07-24).
# Sources: https://docs.greenpt.ai/pricing and https://docs.greenpt.ai/model-cards
# reasoning_effort: thinking is enabled by default, "none" disables it, and
# minimal/low/medium/high are accepted. See https://docs.greenpt.ai/chat-completion
# reasoning_effort: thinking is enabled by default and "none" disables it.
# Values verified against the production API. See https://docs.greenpt.ai/chat-completion
base_model = "google/gemma-4-26b-a4b-it"
name = "gemma4"
reasoning_options = [{ type = "effort", values = ["none", "minimal", "low", "medium", "high"] }]
Expand Down
1 change: 1 addition & 0 deletions providers/greenpt/models/glm-5.1.toml
Original file line number Diff line number Diff line change
@@ -1,6 +1,7 @@
# Cost: converted from GreenPT's EUR list price at 1.14 USD/EUR (rate captured 2026-07-24).
# Sources: https://docs.greenpt.ai/pricing and https://docs.greenpt.ai/model-cards
base_model = "zhipuai/glm-5.1"
status = "deprecated"
reasoning_options = []

[cost]
Expand Down
15 changes: 15 additions & 0 deletions providers/greenpt/models/glm-5.2-caveman-lite.toml
Original file line number Diff line number Diff line change
@@ -0,0 +1,15 @@
# Cost: converted from GreenPT's EUR list price at 1.14 USD/EUR (rate captured 2026-08-01).
# Sources: https://docs.greenpt.ai/pricing and https://docs.greenpt.ai/model-cards
# reasoning_effort: thinking is enabled by default and "none" disables it.
# Values verified against the production API. See https://docs.greenpt.ai/chat-completion
# cache_read: discounted rate for cached prompt tokens. Cache writes are not
# charged. See https://docs.greenpt.ai/prompt-caching
base_model = "zhipuai/glm-5.2"
name = "GLM-5.2 Caveman Lite"
description = "glm-5.2 carrying a built-in ruleset that compresses prose, keeping code and technical detail verbatim. The gentlest tier: it cuts filler only and keeps the explanation intact. Same upstream model and price per token as glm-5.2, with fewer output tokens."
reasoning_options = [{ type = "effort", values = ["none", "minimal", "low", "medium", "high"] }]

[cost]
input = 1.254
cache_read = 0.3135
output = 5.016
15 changes: 15 additions & 0 deletions providers/greenpt/models/glm-5.2-caveman-ultra.toml
Original file line number Diff line number Diff line change
@@ -0,0 +1,15 @@
# Cost: converted from GreenPT's EUR list price at 1.14 USD/EUR (rate captured 2026-08-01).
# Sources: https://docs.greenpt.ai/pricing and https://docs.greenpt.ai/model-cards
# reasoning_effort: thinking is enabled by default and "none" disables it.
# Values verified against the production API. See https://docs.greenpt.ai/chat-completion
# cache_read: discounted rate for cached prompt tokens. Cache writes are not
# charged. See https://docs.greenpt.ai/prompt-caching
base_model = "zhipuai/glm-5.2"
name = "GLM-5.2 Caveman Ultra"
description = "glm-5.2 carrying a built-in ruleset that compresses prose, keeping code and technical detail verbatim. The most aggressive tier, close to answer-only. Same upstream model and price per token as glm-5.2, with fewer output tokens."
reasoning_options = [{ type = "effort", values = ["none", "minimal", "low", "medium", "high"] }]

[cost]
input = 1.254
cache_read = 0.3135
output = 5.016
15 changes: 15 additions & 0 deletions providers/greenpt/models/glm-5.2-caveman.toml
Original file line number Diff line number Diff line change
@@ -0,0 +1,15 @@
# Cost: converted from GreenPT's EUR list price at 1.14 USD/EUR (rate captured 2026-08-01).
# Sources: https://docs.greenpt.ai/pricing and https://docs.greenpt.ai/model-cards
# reasoning_effort: thinking is enabled by default and "none" disables it.
# Values verified against the production API. See https://docs.greenpt.ai/chat-completion
# cache_read: discounted rate for cached prompt tokens. Cache writes are not
# charged. See https://docs.greenpt.ai/prompt-caching
base_model = "zhipuai/glm-5.2"
name = "GLM-5.2 Caveman"
description = "glm-5.2 carrying a built-in ruleset that compresses prose, keeping code and technical detail verbatim. The middle tier, and the ruleset as its authors wrote it. Same upstream model and price per token as glm-5.2, with fewer output tokens."
reasoning_options = [{ type = "effort", values = ["none", "minimal", "low", "medium", "high"] }]

[cost]
input = 1.254
cache_read = 0.3135
output = 5.016
15 changes: 15 additions & 0 deletions providers/greenpt/models/glm-5.2-honey-lite.toml
Original file line number Diff line number Diff line change
@@ -0,0 +1,15 @@
# Cost: converted from GreenPT's EUR list price at 1.14 USD/EUR (rate captured 2026-08-01).
# Sources: https://docs.greenpt.ai/pricing and https://docs.greenpt.ai/model-cards
# reasoning_effort: thinking is enabled by default and "none" disables it.
# Values verified against the production API. See https://docs.greenpt.ai/chat-completion
# cache_read: discounted rate for cached prompt tokens. Cache writes are not
# charged. See https://docs.greenpt.ai/prompt-caching
base_model = "zhipuai/glm-5.2"
name = "GLM-5.2 Honey Lite"
description = "glm-5.2 carrying a built-in ruleset that compresses both generated code and prose. The gentlest tier: it cuts filler only and keeps the explanation intact. Same upstream model and price per token as glm-5.2, with fewer output tokens."
reasoning_options = [{ type = "effort", values = ["none", "minimal", "low", "medium", "high"] }]

[cost]
input = 1.254
cache_read = 0.3135
output = 5.016
15 changes: 15 additions & 0 deletions providers/greenpt/models/glm-5.2-honey-ultra.toml
Original file line number Diff line number Diff line change
@@ -0,0 +1,15 @@
# Cost: converted from GreenPT's EUR list price at 1.14 USD/EUR (rate captured 2026-08-01).
# Sources: https://docs.greenpt.ai/pricing and https://docs.greenpt.ai/model-cards
# reasoning_effort: thinking is enabled by default and "none" disables it.
# Values verified against the production API. See https://docs.greenpt.ai/chat-completion
# cache_read: discounted rate for cached prompt tokens. Cache writes are not
# charged. See https://docs.greenpt.ai/prompt-caching
base_model = "zhipuai/glm-5.2"
name = "GLM-5.2 Honey Ultra"
description = "glm-5.2 carrying a built-in ruleset that compresses both generated code and prose. The most aggressive tier, close to answer-only. Same upstream model and price per token as glm-5.2, with fewer output tokens."
reasoning_options = [{ type = "effort", values = ["none", "minimal", "low", "medium", "high"] }]

[cost]
input = 1.254
cache_read = 0.3135
output = 5.016
15 changes: 15 additions & 0 deletions providers/greenpt/models/glm-5.2-honey.toml
Original file line number Diff line number Diff line change
@@ -0,0 +1,15 @@
# Cost: converted from GreenPT's EUR list price at 1.14 USD/EUR (rate captured 2026-08-01).
# Sources: https://docs.greenpt.ai/pricing and https://docs.greenpt.ai/model-cards
# reasoning_effort: thinking is enabled by default and "none" disables it.
# Values verified against the production API. See https://docs.greenpt.ai/chat-completion
# cache_read: discounted rate for cached prompt tokens. Cache writes are not
# charged. See https://docs.greenpt.ai/prompt-caching
base_model = "zhipuai/glm-5.2"
name = "GLM-5.2 Honey"
description = "glm-5.2 carrying a built-in ruleset that compresses both generated code and prose. The middle tier, and the ruleset as its authors wrote it. Same upstream model and price per token as glm-5.2, with fewer output tokens."
reasoning_options = [{ type = "effort", values = ["none", "minimal", "low", "medium", "high"] }]

[cost]
input = 1.254
cache_read = 0.3135
output = 5.016
15 changes: 15 additions & 0 deletions providers/greenpt/models/glm-5.2-ponytail-lite.toml
Original file line number Diff line number Diff line change
@@ -0,0 +1,15 @@
# Cost: converted from GreenPT's EUR list price at 1.14 USD/EUR (rate captured 2026-08-01).
# Sources: https://docs.greenpt.ai/pricing and https://docs.greenpt.ai/model-cards
# reasoning_effort: thinking is enabled by default and "none" disables it.
# Values verified against the production API. See https://docs.greenpt.ai/chat-completion
# cache_read: discounted rate for cached prompt tokens. Cache writes are not
# charged. See https://docs.greenpt.ai/prompt-caching
base_model = "zhipuai/glm-5.2"
name = "GLM-5.2 Ponytail Lite"
description = "glm-5.2 carrying a built-in ruleset that compresses generated code, preferring platform features over custom code. The gentlest tier: it cuts filler only and keeps the explanation intact. Same upstream model and price per token as glm-5.2, with fewer output tokens."
reasoning_options = [{ type = "effort", values = ["none", "minimal", "low", "medium", "high"] }]

[cost]
input = 1.254
cache_read = 0.3135
output = 5.016
15 changes: 15 additions & 0 deletions providers/greenpt/models/glm-5.2-ponytail-ultra.toml
Original file line number Diff line number Diff line change
@@ -0,0 +1,15 @@
# Cost: converted from GreenPT's EUR list price at 1.14 USD/EUR (rate captured 2026-08-01).
# Sources: https://docs.greenpt.ai/pricing and https://docs.greenpt.ai/model-cards
# reasoning_effort: thinking is enabled by default and "none" disables it.
# Values verified against the production API. See https://docs.greenpt.ai/chat-completion
# cache_read: discounted rate for cached prompt tokens. Cache writes are not
# charged. See https://docs.greenpt.ai/prompt-caching
base_model = "zhipuai/glm-5.2"
name = "GLM-5.2 Ponytail Ultra"
description = "glm-5.2 carrying a built-in ruleset that compresses generated code, preferring platform features over custom code. The most aggressive tier, close to answer-only. Same upstream model and price per token as glm-5.2, with fewer output tokens."
reasoning_options = [{ type = "effort", values = ["none", "minimal", "low", "medium", "high"] }]

[cost]
input = 1.254
cache_read = 0.3135
output = 5.016
15 changes: 15 additions & 0 deletions providers/greenpt/models/glm-5.2-ponytail.toml
Original file line number Diff line number Diff line change
@@ -0,0 +1,15 @@
# Cost: converted from GreenPT's EUR list price at 1.14 USD/EUR (rate captured 2026-08-01).
# Sources: https://docs.greenpt.ai/pricing and https://docs.greenpt.ai/model-cards
# reasoning_effort: thinking is enabled by default and "none" disables it.
# Values verified against the production API. See https://docs.greenpt.ai/chat-completion
# cache_read: discounted rate for cached prompt tokens. Cache writes are not
# charged. See https://docs.greenpt.ai/prompt-caching
base_model = "zhipuai/glm-5.2"
name = "GLM-5.2 Ponytail"
description = "glm-5.2 carrying a built-in ruleset that compresses generated code, preferring platform features over custom code. The middle tier, and the ruleset as its authors wrote it. Same upstream model and price per token as glm-5.2, with fewer output tokens."
reasoning_options = [{ type = "effort", values = ["none", "minimal", "low", "medium", "high"] }]

[cost]
input = 1.254
cache_read = 0.3135
output = 5.016
9 changes: 7 additions & 2 deletions providers/greenpt/models/glm-5.2.toml
Original file line number Diff line number Diff line change
@@ -1,8 +1,13 @@
# Cost: converted from GreenPT's EUR list price at 1.14 USD/EUR (rate captured 2026-07-24).
# Cost: converted from GreenPT's EUR list price at 1.14 USD/EUR (rate captured 2026-08-01).
# Sources: https://docs.greenpt.ai/pricing and https://docs.greenpt.ai/model-cards
# reasoning_effort: thinking is enabled by default and "none" disables it.
# Values verified against the production API. See https://docs.greenpt.ai/chat-completion
# cache_read: discounted rate for cached prompt tokens. Cache writes are not
# charged. See https://docs.greenpt.ai/prompt-caching
base_model = "zhipuai/glm-5.2"
reasoning_options = []
reasoning_options = [{ type = "effort", values = ["none", "minimal", "low", "medium", "high"] }]

[cost]
input = 1.254
cache_read = 0.3135
output = 5.016
5 changes: 4 additions & 1 deletion providers/greenpt/models/gpt-oss-120b.toml
Original file line number Diff line number Diff line change
@@ -1,8 +1,11 @@
# Cost: converted from GreenPT's EUR list price at 1.14 USD/EUR (rate captured 2026-07-24).
# Sources: https://docs.greenpt.ai/pricing and https://docs.greenpt.ai/model-cards
# reasoning_effort: this endpoint accepts only the values below; the remaining
# documented values are rejected with a 400. Verified against the production API.
# See https://docs.greenpt.ai/chat-completion
base_model = "openai/gpt-oss-120b"
attachment = true
reasoning_options = []
reasoning_options = [{ type = "effort", values = ["low", "medium", "high"] }]

[cost]
input = 0.228
Expand Down
7 changes: 4 additions & 3 deletions providers/greenpt/models/green-r-raw.toml
Original file line number Diff line number Diff line change
@@ -1,11 +1,12 @@
# Cost: converted from GreenPT's EUR list price at 1.14 USD/EUR (rate captured 2026-07-24).
# Sources: https://docs.greenpt.ai/pricing and https://docs.greenpt.ai/model-cards
# reasoning_effort: thinking is enabled by default, "none" disables it, and
# minimal/low/medium/high are accepted. See https://docs.greenpt.ai/chat-completion
# reasoning_effort: this endpoint accepts only the values below; the remaining
# documented values are rejected with a 400. Verified against the production API.
# See https://docs.greenpt.ai/chat-completion
base_model = "openai/gpt-oss-120b"
name = "Green R Raw"
attachment = true
reasoning_options = [{ type = "effort", values = ["none", "minimal", "low", "medium", "high"] }]
reasoning_options = [{ type = "effort", values = ["low", "medium", "high"] }]

[cost]
input = 0.399
Expand Down
7 changes: 4 additions & 3 deletions providers/greenpt/models/green-r.toml
Original file line number Diff line number Diff line change
@@ -1,11 +1,12 @@
# Cost: converted from GreenPT's EUR list price at 1.14 USD/EUR (rate captured 2026-07-24).
# Sources: https://docs.greenpt.ai/pricing and https://docs.greenpt.ai/model-cards
# reasoning_effort: thinking is enabled by default, "none" disables it, and
# minimal/low/medium/high are accepted. See https://docs.greenpt.ai/chat-completion
# reasoning_effort: this endpoint accepts only the values below; the remaining
# documented values are rejected with a 400. Verified against the production API.
# See https://docs.greenpt.ai/chat-completion
base_model = "openai/gpt-oss-120b"
name = "Green R"
attachment = true
reasoning_options = [{ type = "effort", values = ["none", "minimal", "low", "medium", "high"] }]
reasoning_options = [{ type = "effort", values = ["low", "medium", "high"] }]

[cost]
input = 0.399
Expand Down
2 changes: 1 addition & 1 deletion providers/greenpt/models/holo2-30b-a3b.toml
Original file line number Diff line number Diff line change
Expand Up @@ -6,7 +6,7 @@ release_date = "2025-11"
last_updated = "2025-11"
attachment = true
reasoning = true
reasoning_options = []
reasoning_options = [{ type = "effort", values = ["low", "medium", "high"] }]
temperature = true
tool_call = true
structured_output = false
Expand Down
1 change: 1 addition & 0 deletions providers/greenpt/models/kimi-k2.6-fast.toml
Original file line number Diff line number Diff line change
Expand Up @@ -2,6 +2,7 @@
# Sources: https://docs.greenpt.ai/pricing and https://docs.greenpt.ai/model-cards
base_model = "moonshotai/kimi-k2.6"
name = "Kimi K2.6 Fast"
status = "deprecated"
reasoning_options = []

[cost]
Expand Down
13 changes: 9 additions & 4 deletions providers/greenpt/models/kimi-k2.6.toml
Original file line number Diff line number Diff line change
@@ -1,11 +1,16 @@
# Cost: converted from GreenPT's EUR list price at 1.14 USD/EUR (rate captured 2026-07-24).
# Cost: converted from GreenPT's EUR list price at 1.14 USD/EUR (rate captured 2026-08-01).
# Sources: https://docs.greenpt.ai/pricing and https://docs.greenpt.ai/model-cards
# reasoning_effort: thinking is enabled by default and "none" disables it.
# Values verified against the production API. See https://docs.greenpt.ai/chat-completion
# cache_read: discounted rate for cached prompt tokens. Cache writes are not
# charged. See https://docs.greenpt.ai/prompt-caching
base_model = "moonshotai/kimi-k2.6"
reasoning_options = []
reasoning_options = [{ type = "effort", values = ["none", "minimal", "low", "medium", "high"] }]

[cost]
input = 0.828
output = 4.389
input = 0.7524
cache_read = 0.2508
output = 4.275

[modalities]
input = ["text", "image"]
11 changes: 8 additions & 3 deletions providers/greenpt/models/kimi-k2.7-code.toml
Original file line number Diff line number Diff line change
@@ -1,11 +1,16 @@
# Cost: converted from GreenPT's EUR list price at 1.14 USD/EUR (rate captured 2026-07-24).
# Cost: converted from GreenPT's EUR list price at 1.14 USD/EUR (rate captured 2026-08-01).
# Sources: https://docs.greenpt.ai/pricing and https://docs.greenpt.ai/model-cards
# reasoning_effort: thinking is enabled by default and "none" disables it.
# Values verified against the production API. See https://docs.greenpt.ai/chat-completion
# cache_read: discounted rate for cached prompt tokens. Cache writes are not
# charged. See https://docs.greenpt.ai/prompt-caching
base_model = "moonshotai/kimi-k2.7-code"
attachment = false
reasoning_options = []
reasoning_options = [{ type = "effort", values = ["none", "minimal", "low", "medium", "high"] }]

[cost]
input = 0.941
input = 0.9006
cache_read = 0.1881
output = 4.389

[modalities]
Expand Down
16 changes: 16 additions & 0 deletions providers/greenpt/models/kimi-k3.toml
Original file line number Diff line number Diff line change
@@ -0,0 +1,16 @@
# Cost: converted from GreenPT's EUR list price at 1.14 USD/EUR (rate captured 2026-08-01).
# Sources: https://docs.greenpt.ai/pricing and https://docs.greenpt.ai/model-cards
# reasoning_effort: thinking is enabled by default and "none" disables it.
# Values verified against the production API. See https://docs.greenpt.ai/chat-completion
# cache_read: discounted rate for cached prompt tokens. Cache writes are not
# charged. See https://docs.greenpt.ai/prompt-caching
base_model = "moonshotai/kimi-k3"
reasoning_options = [{ type = "effort", values = ["none", "minimal", "low", "medium", "high"] }]

[cost]
input = 3.762
cache_read = 0.9405
output = 18.81

[modalities]
input = ["text", "image"]
11 changes: 8 additions & 3 deletions providers/greenpt/models/minimax-m2.5.toml
Original file line number Diff line number Diff line change
@@ -1,8 +1,13 @@
# Cost: converted from GreenPT's EUR list price at 1.14 USD/EUR (rate captured 2026-07-24).
# Cost: converted from GreenPT's EUR list price at 1.14 USD/EUR (rate captured 2026-08-01).
# Sources: https://docs.greenpt.ai/pricing and https://docs.greenpt.ai/model-cards
# reasoning_effort: thinking is enabled by default and "none" disables it.
# Values verified against the production API. See https://docs.greenpt.ai/chat-completion
# cache_read: discounted rate for cached prompt tokens. Cache writes are not
# charged. See https://docs.greenpt.ai/prompt-caching
base_model = "minimax/MiniMax-M2.5"
reasoning_options = []
reasoning_options = [{ type = "effort", values = ["none", "minimal", "low", "medium", "high"] }]

[cost]
input = 0.188
input = 0.1938
cache_read = 0.0627
output = 1.129
5 changes: 4 additions & 1 deletion providers/greenpt/models/mistral-medium-3.5-128b.toml
Original file line number Diff line number Diff line change
@@ -1,7 +1,10 @@
# Cost: converted from GreenPT's EUR list price at 1.14 USD/EUR (rate captured 2026-07-24).
# Sources: https://docs.greenpt.ai/pricing and https://docs.greenpt.ai/model-cards
# reasoning_effort: this endpoint accepts only the values below; the remaining
# documented values are rejected with a 400. Verified against the production API.
# See https://docs.greenpt.ai/chat-completion
base_model = "mistral/mistral-medium-2604"
reasoning_options = []
reasoning_options = [{ type = "effort", values = ["none", "high"] }]

[cost]
input = 2.052
Expand Down
4 changes: 3 additions & 1 deletion providers/greenpt/models/qwen3.5-397b-a17b.toml
Original file line number Diff line number Diff line change
@@ -1,8 +1,10 @@
# Cost: converted from GreenPT's EUR list price at 1.14 USD/EUR (rate captured 2026-07-24).
# Sources: https://docs.greenpt.ai/pricing and https://docs.greenpt.ai/model-cards
# reasoning_effort: thinking is enabled by default and "none" disables it.
# Values verified against the production API. See https://docs.greenpt.ai/chat-completion
base_model = "alibaba/qwen3.5-397b-a17b"
attachment = false
reasoning_options = []
reasoning_options = [{ type = "effort", values = ["none", "minimal", "low", "medium", "high"] }]

[cost]
input = 0.798
Expand Down
4 changes: 3 additions & 1 deletion providers/greenpt/models/qwen3.6-35b-a3b.toml
Original file line number Diff line number Diff line change
@@ -1,7 +1,9 @@
# Cost: converted from GreenPT's EUR list price at 1.14 USD/EUR (rate captured 2026-07-24).
# Sources: https://docs.greenpt.ai/pricing and https://docs.greenpt.ai/model-cards
# reasoning_effort: thinking is enabled by default and "none" disables it.
# Values verified against the production API. See https://docs.greenpt.ai/chat-completion
base_model = "alibaba/qwen3.6-35b-a3b"
reasoning_options = []
reasoning_options = [{ type = "effort", values = ["none", "minimal", "low", "medium", "high"] }]

[cost]
input = 0.342
Expand Down
Loading