Skip to content
1 change: 1 addition & 0 deletions providers/gptzzz/logo.svg
Loading
Sorry, something went wrong. Reload?
Sorry, we cannot display this file.
Sorry, this file is invalid so it cannot be displayed.
24 changes: 24 additions & 0 deletions providers/gptzzz/models/gpt-5.6-sol.toml
Original file line number Diff line number Diff line change
@@ -0,0 +1,24 @@
# GPTZZZ is an OpenAI-compatible relay. Effort values verified per model against the
# live /v1/chat/completions endpoint on 2026-09-16 and re-probed on 2026-09-29: none, low,
# medium, high, xhigh and max all return HTTP 200; minimal returns HTTP 400.
# Pricing source: https://gptzzz.ai/model-plaza (read from /api/v1/model-plaza on 2026-09-29).
# GPTZZZ bills this model against its own published list price of 5.00 input / 30.00 output /
# 0.50 cache read / 6.25 cache write per MTok (>272K context: 10.00 / 45.00 / 1.00 / 12.50),
# which is higher than OpenAI's list for gpt-5.6-sol, multiplied by its public 0.05x group rate.
# The host's /v1/usage "cost" field matches this 5.00 / 30.00 list exactly, so the values
# below are what the host actually charges.
base_model = "openai/gpt-5.6-sol"
reasoning_options = [{ type = "effort", values = ["none", "low", "medium", "high", "xhigh", "max"] }]

[cost]
input = 0.25
output = 1.5
cache_read = 0.025
cache_write = 0.3125

[[cost.tiers]]
tier = { size = 272_000 }
input = 0.5
output = 2.25
cache_read = 0.05
cache_write = 0.625
21 changes: 21 additions & 0 deletions providers/gptzzz/models/gpt-5.6-terra.toml
Original file line number Diff line number Diff line change
@@ -0,0 +1,21 @@
# GPTZZZ is an OpenAI-compatible relay. Effort values verified per model against the
# live /v1/chat/completions endpoint on 2026-09-16 and re-probed on 2026-09-29: none, low,
# medium, high, xhigh and max all return HTTP 200; minimal returns HTTP 400.
# Pricing source: https://gptzzz.ai/model-plaza (read from /api/v1/model-plaza on 2026-09-29).
# Published list price 2.00 / 12.00 / 0.20 / 2.50 per MTok (>272K: 4.00 / 18.00 / 0.40 / 5.00),
# multiplied by the public 0.05x group rate. The host's /v1/usage "cost" field matches this list.
base_model = "openai/gpt-5.6-terra"
reasoning_options = [{ type = "effort", values = ["none", "low", "medium", "high", "xhigh", "max"] }]

[cost]
input = 0.1
output = 0.6
cache_read = 0.01
cache_write = 0.125

[[cost.tiers]]
tier = { size = 272_000 }
input = 0.2
output = 0.9
cache_read = 0.02
cache_write = 0.25
25 changes: 25 additions & 0 deletions providers/gptzzz/models/gpt-5.6.toml
Original file line number Diff line number Diff line change
@@ -0,0 +1,25 @@
# GPTZZZ is an OpenAI-compatible relay. Effort values verified per model against the
# live /v1/chat/completions endpoint on 2026-09-16 and re-probed on 2026-09-29: none, low,
# medium, high, xhigh and max all return HTTP 200; minimal returns HTTP 400.
# Pricing source: https://gptzzz.ai/model-plaza (read from /api/v1/model-plaza on 2026-09-29).
# GPTZZZ bills this model against its own published list price of 5.00 input / 30.00 output /
# 0.50 cache read / 6.25 cache write per MTok (>272K context: 10.00 / 45.00 / 1.00 / 12.50),
# which is higher than OpenAI's list for gpt-5.6-sol, multiplied by its public 0.05x group rate.
# The host's /v1/usage "cost" field matches this 5.00 / 30.00 list exactly, so the values
# below are what the host actually charges.
base_model = "openai/gpt-5.6-sol"
name = "GPT-5.6"
reasoning_options = [{ type = "effort", values = ["none", "low", "medium", "high", "xhigh", "max"] }]

[cost]
input = 0.25
output = 1.5
cache_read = 0.025
cache_write = 0.3125

[[cost.tiers]]
tier = { size = 272_000 }
input = 0.5
output = 2.25
cache_read = 0.05
cache_write = 0.625
24 changes: 24 additions & 0 deletions providers/gptzzz/models/gpt-6-astra.toml
Original file line number Diff line number Diff line change
@@ -0,0 +1,24 @@
# GPTZZZ is an OpenAI-compatible relay. Effort values probed one by one against the live
# /v1/chat/completions endpoint on 2026-09-29: low, medium, high, xhigh and max return
# HTTP 200 (reasoning_tokens 10 / 11 / 14 at high / xhigh / max); none and minimal return
# HTTP 400 ("Supported values are: 'low', 'medium', 'high', 'xhigh', and 'max'").
# This matches the OpenAI lab entry.
# Pricing source: https://gptzzz.ai/model-plaza (read from /api/v1/model-plaza on 2026-09-29).
# Published list price 10.00 / 50.00 / 1.00 / 12.50 per MTok (>272K: 20.00 / 75.00 / 2.00 / 25.00),
# multiplied by the public 0.05x group rate. The host's /v1/usage "cost" field matches this
# list price exactly.
base_model = "openai/gpt-6-astra"
reasoning_options = [{ type = "effort", values = ["low", "medium", "high", "xhigh", "max"] }]

[cost]
input = 0.5
output = 2.5
cache_read = 0.05
cache_write = 0.625

[[cost.tiers]]
tier = { size = 272_000 }
input = 1.0
output = 3.75
cache_read = 0.1
cache_write = 1.25
5 changes: 5 additions & 0 deletions providers/gptzzz/provider.toml
Original file line number Diff line number Diff line change
@@ -0,0 +1,5 @@
name = "GPTZZZ"
env = ["GPTZZZ_API_KEY"]
npm = "@ai-sdk/openai-compatible"
api = "https://gptzzz.ai/v1"
doc = "https://gptzzz.ai/docs/"
Loading