Skip to content
Merged
Show file tree
Hide file tree
Changes from all commits
Commits
File filter

Filter by extension

Filter by extension

Conversations
Failed to load comments.
Loading
Jump to
Jump to file
Failed to load files.
Loading
Diff view
Diff view
6 changes: 6 additions & 0 deletions .github/scripts/autogen/types.ts
Original file line number Diff line number Diff line change
Expand Up @@ -157,6 +157,12 @@ export interface components {
supportedModes?: components["schemas"]["Mode"][];
/** @description Whether the model supports extended thinking / reasoning */
thinking?: boolean;
/**
* @description When true, inbound Chat Completions stay on /chat/completions unless a
* responses-only model or an existing trigger (header / reasoning / Cursor).
* Omit on new models so Chat Completions is always served via Responses.
*/
preferChatCompletions?: boolean;
};
ModelParam: {
defaultValue?: (string | number | boolean) | null;
Expand Down
4 changes: 4 additions & 0 deletions .github/test/model.cue
Original file line number Diff line number Diff line change
Expand Up @@ -341,6 +341,10 @@ package model
supportedModes?: [...#Mode]
// Whether the model supports extended thinking / reasoning
thinking?: bool
// When true, inbound Chat Completions stay on /chat/completions unless a
// responses-only model or an existing trigger (header / reasoning / Cursor).
// Omit on new models so Chat Completions is always served via Responses.
preferChatCompletions?: bool
}

#PricingTier: {
Expand Down
8 changes: 8 additions & 0 deletions CONTRIBUTING.md
Original file line number Diff line number Diff line change
Expand Up @@ -147,6 +147,14 @@ removeParams: [<param-key>, ...]
# Param keys that must always be provided by callers
requiredParams: [<param-key>, ...]

# Whether the model supports extended thinking / reasoning
thinking: true

# When true, inbound Chat Completions stay on /chat/completions unless a
# responses-only model or an existing trigger (header / reasoning / Cursor).
# Omit on new models so Chat Completions is always served via Responses.
preferChatCompletions: true

# Provider-specific escape-hatch configuration. Only set a key here when the
# behaviour cannot be expressed through the generic fields above, and only on
# models belonging to the provider that owns the key.
Expand Down
1 change: 1 addition & 0 deletions providers/aws-bedrock-mantle/google.gemma-4-26b-a4b.yaml
Original file line number Diff line number Diff line change
Expand Up @@ -30,6 +30,7 @@ modalities:
- text
mode: chat
model: google.gemma-4-26b-a4b
preferChatCompletions: true
provisioning: serverless
sources:
- https://docs.aws.amazon.com/bedrock/latest/userguide/model-card-google-gemma-4-26b-a4b.html
Expand Down
1 change: 1 addition & 0 deletions providers/aws-bedrock-mantle/google.gemma-4-31b.yaml
Original file line number Diff line number Diff line change
Expand Up @@ -27,6 +27,7 @@ modalities:
- text
mode: chat
model: google.gemma-4-31b
preferChatCompletions: true
provisioning: serverless
sources:
- https://docs.aws.amazon.com/bedrock/latest/userguide/model-card-google-gemma-4-31b.html
Expand Down
1 change: 1 addition & 0 deletions providers/aws-bedrock-mantle/google.gemma-4-e2b.yaml
Original file line number Diff line number Diff line change
Expand Up @@ -28,6 +28,7 @@ modalities:
- text
mode: chat
model: google.gemma-4-e2b
preferChatCompletions: true
provisioning: serverless
sources:
- https://docs.aws.amazon.com/bedrock/latest/userguide/model-card-google-gemma-4-e2b.html
Expand Down
1 change: 1 addition & 0 deletions providers/aws-bedrock-mantle/openai.gpt-oss-120b.yaml
Original file line number Diff line number Diff line change
Expand Up @@ -69,6 +69,7 @@ params:
- medium
- high
type: string
preferChatCompletions: true
provisioning: serverless
sources:
- https://docs.aws.amazon.com/bedrock/latest/userguide/model-card-openai-gpt-oss-120b.html
Expand Down
1 change: 1 addition & 0 deletions providers/aws-bedrock-mantle/openai.gpt-oss-20b.yaml
Original file line number Diff line number Diff line change
Expand Up @@ -61,6 +61,7 @@ params:
- medium
- high
type: string
preferChatCompletions: true
provisioning: serverless
sources:
- https://docs.aws.amazon.com/bedrock/latest/userguide/model-card-openai-gpt-oss-20b.html
Expand Down
1 change: 1 addition & 0 deletions providers/aws-bedrock-mantle/xai.grok-4.3.yaml
Original file line number Diff line number Diff line change
Expand Up @@ -51,6 +51,7 @@ params:
- defaultValue: 131072
key: max_completion_tokens
type: number
preferChatCompletions: true
provisioning: serverless
sources:
- https://docs.aws.amazon.com/bedrock/latest/userguide/model-card-xai-grok-4-3.html
Expand Down
1 change: 1 addition & 0 deletions providers/aws-bedrock-mantle/xai.grok-4.6.yaml
Original file line number Diff line number Diff line change
Expand Up @@ -39,6 +39,7 @@ params:
- defaultValue: 131072
key: max_completion_tokens
type: number
preferChatCompletions: true
provisioning: serverless
sources:
- https://docs.aws.amazon.com/bedrock/latest/userguide/model-card-xai-grok-4-6.html
Expand Down
1 change: 0 additions & 1 deletion providers/aws-bedrock/openai.gpt-oss-120b-1:0.yaml
Original file line number Diff line number Diff line change
Expand Up @@ -84,5 +84,4 @@ sources:
status: active
supportedModes:
- chat
- responses
thinking: true
1 change: 0 additions & 1 deletion providers/aws-bedrock/openai.gpt-oss-20b-1:0.yaml
Original file line number Diff line number Diff line change
Expand Up @@ -85,5 +85,4 @@ sources:
status: active
supportedModes:
- chat
- responses
thinking: true
Original file line number Diff line number Diff line change
Expand Up @@ -27,6 +27,7 @@ params:
key: max_tokens
maxValue: 1024
minValue: 1
preferChatCompletions: true
provisioning: serverless
status: deprecated
supportedModes:
Expand Down
1 change: 1 addition & 0 deletions providers/azure-open-ai/gpt-4.1-2025-04-14.yaml
Original file line number Diff line number Diff line change
Expand Up @@ -35,6 +35,7 @@ model: gpt-4.1-2025-04-14
params:
- key: max_tokens
maxValue: 32768
preferChatCompletions: true
retirementDate: "2027-04-14"
sources:
- https://learn.microsoft.com/en-us/azure/ai-foundry/openai/concepts/models
Expand Down
1 change: 1 addition & 0 deletions providers/azure-open-ai/gpt-4.1-mini-2025-04-14.yaml
Original file line number Diff line number Diff line change
Expand Up @@ -35,6 +35,7 @@ model: gpt-4.1-mini-2025-04-14
params:
- key: max_tokens
maxValue: 32768
preferChatCompletions: true
provisioning: serverless
retirementDate: "2027-04-14"
sources:
Expand Down
1 change: 1 addition & 0 deletions providers/azure-open-ai/gpt-4.1-mini.yaml
Original file line number Diff line number Diff line change
Expand Up @@ -37,6 +37,7 @@ params:
key: max_completion_tokens
maxValue: 32768
minValue: 1
preferChatCompletions: true
provisioning: serverless
removeParams:
- max_tokens
Expand Down
1 change: 1 addition & 0 deletions providers/azure-open-ai/gpt-4.1-nano-2025-04-14.yaml
Original file line number Diff line number Diff line change
Expand Up @@ -37,6 +37,7 @@ params:
key: max_tokens
maxValue: 32768
minValue: 1
preferChatCompletions: true
removeParams:
- reasoning_effort
retirementDate: "2027-04-14"
Expand Down
1 change: 1 addition & 0 deletions providers/azure-open-ai/gpt-4.1-nano.yaml
Original file line number Diff line number Diff line change
Expand Up @@ -40,6 +40,7 @@ params:
- defaultValue: null
key: response_format
type: string
preferChatCompletions: true
provisioning: serverless
removeParams:
- max_tokens
Expand Down
1 change: 1 addition & 0 deletions providers/azure-open-ai/gpt-4.1.yaml
Original file line number Diff line number Diff line change
Expand Up @@ -46,6 +46,7 @@ params:
- defaultValue: null
key: response_format
type: string
preferChatCompletions: true
provisioning: serverless
removeParams:
- max_tokens
Expand Down
1 change: 1 addition & 0 deletions providers/azure-open-ai/gpt-4o-2024-05-13.yaml
Original file line number Diff line number Diff line change
Expand Up @@ -23,6 +23,7 @@ params:
- defaultValue: null
key: response_format
type: string
preferChatCompletions: true
supportedModes:
- chat
- responses
1 change: 1 addition & 0 deletions providers/azure-open-ai/gpt-4o-2024-08-06.yaml
Original file line number Diff line number Diff line change
Expand Up @@ -33,6 +33,7 @@ params:
key: max_tokens
maxValue: 16384
minValue: 1
preferChatCompletions: true
sources:
- https://learn.microsoft.com/en-us/azure/ai-foundry/openai/concepts/models
status: deprecated
Expand Down
1 change: 1 addition & 0 deletions providers/azure-open-ai/gpt-4o-2024-11-20.yaml
Original file line number Diff line number Diff line change
Expand Up @@ -39,6 +39,7 @@ model: gpt-4o-2024-11-20
params:
- key: max_tokens
maxValue: 16384
preferChatCompletions: true
provisioning: serverless
removeParams:
- reasoning_effort
Expand Down
1 change: 1 addition & 0 deletions providers/azure-open-ai/gpt-4o-mini-2024-07-18.yaml
Original file line number Diff line number Diff line change
Expand Up @@ -38,6 +38,7 @@ model: gpt-4o-mini-2024-07-18
params:
- key: max_tokens
maxValue: 16384
preferChatCompletions: true
provisioning: serverless
retirementDate: "2027-04-14"
status: deprecated
Expand Down
1 change: 1 addition & 0 deletions providers/azure-open-ai/gpt-4o-mini.yaml
Original file line number Diff line number Diff line change
Expand Up @@ -40,6 +40,7 @@ params:
key: max_tokens
maxValue: 16384
minValue: 1
preferChatCompletions: true
provisioning: serverless
retirementDate: "2027-04-14"
sources:
Expand Down
1 change: 1 addition & 0 deletions providers/azure-open-ai/gpt-4o.yaml
Original file line number Diff line number Diff line change
Expand Up @@ -35,6 +35,7 @@ model: gpt-4o
params:
- key: max_tokens
maxValue: 16384
preferChatCompletions: true
provisioning: serverless
retirementDate: "2026-10-01"
status: deprecated
Expand Down
1 change: 1 addition & 0 deletions providers/azure-open-ai/gpt-5-2025-08-07.yaml
Original file line number Diff line number Diff line change
Expand Up @@ -49,6 +49,7 @@ params:
- medium
- high
type: string
preferChatCompletions: true
provisioning: serverless
removeParams:
- max_tokens
Expand Down
1 change: 1 addition & 0 deletions providers/azure-open-ai/gpt-5-chat-2025-08-07.yaml
Original file line number Diff line number Diff line change
Expand Up @@ -31,6 +31,7 @@ params:
key: max_completion_tokens
maxValue: 16384
minValue: 1
preferChatCompletions: true
removeParams:
- max_tokens
sources:
Expand Down
1 change: 1 addition & 0 deletions providers/azure-open-ai/gpt-5-chat-2025-10-03.yaml
Original file line number Diff line number Diff line change
Expand Up @@ -28,6 +28,7 @@ model: gpt-5-chat-2025-10-03
params:
- key: max_completion_tokens
maxValue: 16384
preferChatCompletions: true
provisioning: serverless
removeParams:
- max_tokens
Expand Down
1 change: 1 addition & 0 deletions providers/azure-open-ai/gpt-5-chat-latest.yaml
Original file line number Diff line number Diff line change
Expand Up @@ -29,6 +29,7 @@ params:
key: max_completion_tokens
maxValue: 16384
minValue: 1
preferChatCompletions: true
provisioning: serverless
removeParams:
- max_tokens
Expand Down
1 change: 1 addition & 0 deletions providers/azure-open-ai/gpt-5-chat.yaml
Original file line number Diff line number Diff line change
Expand Up @@ -30,6 +30,7 @@ params:
key: max_completion_tokens
maxValue: 16384
minValue: 1
preferChatCompletions: true
provisioning: serverless
removeParams:
- max_tokens
Expand Down
1 change: 1 addition & 0 deletions providers/azure-open-ai/gpt-5-mini-2025-08-07-lite.yaml
Original file line number Diff line number Diff line change
Expand Up @@ -43,6 +43,7 @@ params:
key: max_completion_tokens
maxValue: 128000
minValue: 1
preferChatCompletions: true
provisioning: serverless
removeParams:
- max_tokens
Expand Down
1 change: 1 addition & 0 deletions providers/azure-open-ai/gpt-5-mini-2025-08-07.yaml
Original file line number Diff line number Diff line change
Expand Up @@ -57,6 +57,7 @@ params:
- medium
- high
type: string
preferChatCompletions: true
provisioning: serverless
removeParams:
- max_tokens
Expand Down
1 change: 1 addition & 0 deletions providers/azure-open-ai/gpt-5-mini.yaml
Original file line number Diff line number Diff line change
Expand Up @@ -56,6 +56,7 @@ params:
- medium
- high
type: string
preferChatCompletions: true
provisioning: serverless
removeParams:
- temperature
Expand Down
1 change: 1 addition & 0 deletions providers/azure-open-ai/gpt-5-nano-2025-08-07.yaml
Original file line number Diff line number Diff line change
Expand Up @@ -45,6 +45,7 @@ params:
- medium
- high
type: string
preferChatCompletions: true
provisioning: serverless
removeParams:
- max_tokens
Expand Down
1 change: 1 addition & 0 deletions providers/azure-open-ai/gpt-5-nano.yaml
Original file line number Diff line number Diff line change
Expand Up @@ -58,6 +58,7 @@ params:
- medium
- high
type: string
preferChatCompletions: true
provisioning: serverless
removeParams:
- temperature
Expand Down
1 change: 1 addition & 0 deletions providers/azure-open-ai/gpt-5.1-2025-11-13.yaml
Original file line number Diff line number Diff line change
Expand Up @@ -59,6 +59,7 @@ params:
- medium
- high
type: string
preferChatCompletions: true
provisioning: serverless
removeParams:
- temperature
Expand Down
1 change: 1 addition & 0 deletions providers/azure-open-ai/gpt-5.1-chat-2025-11-13.yaml
Original file line number Diff line number Diff line change
Expand Up @@ -38,6 +38,7 @@ params:
- defaultValue: null
key: reasoning_effort
type: string
preferChatCompletions: true
provisioning: serverless
removeParams:
- temperature
Expand Down
1 change: 1 addition & 0 deletions providers/azure-open-ai/gpt-5.1-chat.yaml
Original file line number Diff line number Diff line change
Expand Up @@ -48,6 +48,7 @@ params:
- high
- xhigh
type: string
preferChatCompletions: true
provisioning: serverless
removeParams:
- temperature
Expand Down
1 change: 1 addition & 0 deletions providers/azure-open-ai/gpt-5.1.yaml
Original file line number Diff line number Diff line change
Expand Up @@ -59,6 +59,7 @@ params:
- high
- xhigh
type: string
preferChatCompletions: true
provisioning: serverless
removeParams:
- temperature
Expand Down
1 change: 1 addition & 0 deletions providers/azure-open-ai/gpt-5.2-2025-12-11.yaml
Original file line number Diff line number Diff line change
Expand Up @@ -43,6 +43,7 @@ params:
- medium
- high
type: string
preferChatCompletions: true
provisioning: serverless
removeParams:
- temperature
Expand Down
1 change: 1 addition & 0 deletions providers/azure-open-ai/gpt-5.2-chat-2025-12-11.yaml
Original file line number Diff line number Diff line change
Expand Up @@ -36,6 +36,7 @@ params:
- medium
- high
type: string
preferChatCompletions: true
provisioning: serverless
removeParams:
- max_tokens
Expand Down
1 change: 1 addition & 0 deletions providers/azure-open-ai/gpt-5.2-chat-2026-02-10.yaml
Original file line number Diff line number Diff line change
Expand Up @@ -28,6 +28,7 @@ model: gpt-5.2-chat-2026-02-10
params:
- key: max_completion_tokens
maxValue: 16384
preferChatCompletions: true
provisioning: serverless
removeParams:
- max_tokens
Expand Down
1 change: 1 addition & 0 deletions providers/azure-open-ai/gpt-5.2-chat.yaml
Original file line number Diff line number Diff line change
Expand Up @@ -28,6 +28,7 @@ model: gpt-5.2-chat
params:
- key: max_completion_tokens
maxValue: 16384
preferChatCompletions: true
provisioning: serverless
removeParams:
- max_tokens
Expand Down
1 change: 1 addition & 0 deletions providers/azure-open-ai/gpt-5.2.yaml
Original file line number Diff line number Diff line change
Expand Up @@ -47,6 +47,7 @@ params:
- defaultValue: null
key: response_format
type: string
preferChatCompletions: true
provisioning: serverless
removeParams:
- temperature
Expand Down
1 change: 1 addition & 0 deletions providers/azure-open-ai/gpt-5.3-chat-2026-03-03.yaml
Original file line number Diff line number Diff line change
Expand Up @@ -30,6 +30,7 @@ params:
key: max_completion_tokens
maxValue: 16384
minValue: 1
preferChatCompletions: true
provisioning: serverless
removeParams:
- max_tokens
Expand Down
1 change: 1 addition & 0 deletions providers/azure-open-ai/gpt-5.yaml
Original file line number Diff line number Diff line change
Expand Up @@ -55,6 +55,7 @@ params:
- medium
- high
type: string
preferChatCompletions: true
provisioning: serverless
removeParams:
- temperature
Expand Down
1 change: 1 addition & 0 deletions providers/azure-open-ai/gpt-chat-latest-2026-05-05.yaml
Original file line number Diff line number Diff line change
Expand Up @@ -23,6 +23,7 @@ modalities:
- text
mode: chat
model: gpt-chat-latest-2026-05-05
preferChatCompletions: true
provisioning: serverless
retirementDate: "2026-08-05"
sources:
Expand Down
1 change: 1 addition & 0 deletions providers/azure-open-ai/gpt-chat-latest-2026-05-28.yaml
Original file line number Diff line number Diff line change
Expand Up @@ -25,6 +25,7 @@ modalities:
- text
mode: chat
model: gpt-chat-latest-2026-05-28
preferChatCompletions: true
provisioning: serverless
retirementDate: "2026-08-28"
status: retired
Expand Down
Loading
Loading