mirror of
https://github.com/supabase/supabase.git
synced 2026-09-22 13:37:53 +08:00
9b17ce8f2c
## I have read the [CONTRIBUTING.md](https://github.com/supabase/supabase/blob/master/CONTRIBUTING.md) file. YES ## What kind of change does this PR introduce? Feature / chore: hide assistant model selection in the UI and default chats to GPT-5.6 Luna. ## What is the current behavior? The assistant composer exposes a model picker. Paid orgs default to `gpt-5.3-codex`; everyone else defaults to `gpt-5.4-nano`. ## What is the new behavior? - The model picker is hidden in the assistant composer and Explorer home. - Chats default to `gpt-5.6-luna` with `reasoningEffort: medium`. - Model selection plumbing is kept (registry, entitlements, `setModel`, generate-v4 request body) so a requested model can still be honored when provided. - Other completion endpoints still use `gpt-5.4-nano`. ## Additional context Model selector UI can be re-enabled by passing `selectedModel` / `onSelectModel` to `AssistantChatForm`. <!-- This is an auto-generated comment: release notes by coderabbit.ai --> ## Summary by CodeRabbit * **New Features** * Added support for the GPT-5.6 Luna model with medium reasoning capability. * Made GPT-5.6 Luna the default assistant model. * **Improvements** * Simplified assistant chat by removing model selection from the primary chat experience. * Updated model fallback behavior to use the standard assistant model. * Chat forms can now optionally display model selection when configured. * **Tests** * Updated model coverage and assistant chat tests for the new defaults and behavior. <!-- end of auto-generated comment: release notes by coderabbit.ai --> --------- Co-authored-by: Cursor <cursoragent@cursor.com>
179 lines
5.7 KiB
TypeScript
179 lines
5.7 KiB
TypeScript
export type ProviderName = 'bedrock' | 'openai'
|
|
|
|
export type BedrockModel = 'anthropic.claude-3-7-sonnet-20250219-v1:0' | 'openai.gpt-oss-120b-1:0'
|
|
|
|
export type OpenAIModelId = 'gpt-5.4-nano' | 'gpt-5.3-codex' | 'gpt-5.6-luna'
|
|
|
|
// Source: https://developers.openai.com/api/docs/guides/reasoning + per-model pages
|
|
export type ReasoningEffort = 'none' | 'minimal' | 'low' | 'medium' | 'high' | 'xhigh' | 'max'
|
|
|
|
// Per-model reasoning effort compatibility.
|
|
// Sources: https://developers.openai.com/api/docs/models/gpt-5.4-nano
|
|
// https://developers.openai.com/api/docs/models/gpt-5.3-codex
|
|
// https://developers.openai.com/api/docs/models/gpt-5.6-luna
|
|
type ModelReasoningSupport = {
|
|
'gpt-5.4-nano': 'none' | 'low' | 'medium' | 'high' | 'xhigh'
|
|
'gpt-5.3-codex': 'low' | 'medium' | 'high' | 'xhigh'
|
|
'gpt-5.6-luna': 'none' | 'low' | 'medium' | 'high' | 'xhigh' | 'max'
|
|
}
|
|
|
|
type ReasoningEffortFor<ModelId extends OpenAIModelId> = ModelId extends keyof ModelReasoningSupport
|
|
? ModelReasoningSupport[ModelId]
|
|
: never
|
|
|
|
/** Type-safe factory for configuring OpenAI models with compatible reasoning efforts. */
|
|
export function openaiModelEntry<
|
|
ModelId extends OpenAIModelId,
|
|
RequiresAdvance extends boolean = false,
|
|
>(config: {
|
|
id: ModelId
|
|
/** When true, the model requires the `assistant.advance_model` entitlement (paid plans). Defaults to false. */
|
|
requiresAdvanceModelEntitlement?: RequiresAdvance
|
|
/**
|
|
* When omitted, OpenAI applies its own default reasoning effort for the model,
|
|
* which may not be zero. Use an explicit level to control cost and latency.
|
|
*/
|
|
reasoningEffort?: ReasoningEffortFor<ModelId>
|
|
}): {
|
|
id: ModelId
|
|
requiresAdvanceModelEntitlement: RequiresAdvance
|
|
reasoningEffort?: ReasoningEffortFor<ModelId>
|
|
} {
|
|
return {
|
|
requiresAdvanceModelEntitlement: false as RequiresAdvance,
|
|
...config,
|
|
}
|
|
}
|
|
|
|
export type OpenAIModelEntry = ReturnType<typeof openaiModelEntry>
|
|
|
|
/** Default model entry for simple completion endpoints where latency is more important than reasoning. */
|
|
export const DEFAULT_COMPLETION_MODEL = openaiModelEntry({
|
|
id: 'gpt-5.4-nano',
|
|
reasoningEffort: 'none',
|
|
})
|
|
|
|
export const LOGS_REWRITE_MODEL = openaiModelEntry({
|
|
id: 'gpt-5.4-nano',
|
|
reasoningEffort: 'low',
|
|
})
|
|
|
|
// Single source of truth for all Assistant chat model variants and their reasoning levels.
|
|
// Models with requiresAdvanceModelEntitlement false are available to all users; true requires the assistant.advance_model entitlement.
|
|
export const ASSISTANT_MODELS = [
|
|
openaiModelEntry({
|
|
id: 'gpt-5.6-luna',
|
|
requiresAdvanceModelEntitlement: false,
|
|
reasoningEffort: 'medium',
|
|
}),
|
|
openaiModelEntry({
|
|
id: 'gpt-5.4-nano',
|
|
requiresAdvanceModelEntitlement: false,
|
|
reasoningEffort: 'low',
|
|
}),
|
|
openaiModelEntry({
|
|
id: 'gpt-5.3-codex',
|
|
requiresAdvanceModelEntitlement: true,
|
|
reasoningEffort: 'low',
|
|
}),
|
|
] as const
|
|
|
|
export type AssistantBaseModelId = Extract<
|
|
(typeof ASSISTANT_MODELS)[number],
|
|
{ requiresAdvanceModelEntitlement: false }
|
|
>['id']
|
|
export type AssistantModelId = (typeof ASSISTANT_MODELS)[number]['id']
|
|
|
|
const ASSISTANT_MODELS_MAP = Object.fromEntries(ASSISTANT_MODELS.map((m) => [m.id, m])) as Record<
|
|
AssistantModelId,
|
|
(typeof ASSISTANT_MODELS)[number]
|
|
>
|
|
|
|
export const DEFAULT_ASSISTANT_BASE_MODEL_ID = 'gpt-5.6-luna' satisfies AssistantBaseModelId
|
|
|
|
export const DEFAULT_ASSISTANT_ADVANCE_MODEL_ID = 'gpt-5.3-codex' satisfies AssistantModelId
|
|
|
|
export function defaultAssistantModelId(_hasAccessToAdvanceModel: boolean): AssistantModelId {
|
|
return DEFAULT_ASSISTANT_BASE_MODEL_ID
|
|
}
|
|
|
|
export function isKnownAssistantModelId(id: string): id is AssistantModelId {
|
|
return Object.hasOwn(ASSISTANT_MODELS_MAP, id)
|
|
}
|
|
|
|
export function isAssistantBaseModelId(id: string): id is AssistantBaseModelId {
|
|
return (
|
|
id in ASSISTANT_MODELS_MAP &&
|
|
!ASSISTANT_MODELS_MAP[id as AssistantModelId].requiresAdvanceModelEntitlement
|
|
)
|
|
}
|
|
|
|
export function isAdvanceOnlyModelId(id: string): boolean {
|
|
return (
|
|
id in ASSISTANT_MODELS_MAP &&
|
|
ASSISTANT_MODELS_MAP[id as AssistantModelId].requiresAdvanceModelEntitlement
|
|
)
|
|
}
|
|
|
|
export function getAssistantModelEntry(id: AssistantModelId): (typeof ASSISTANT_MODELS)[number] {
|
|
return ASSISTANT_MODELS_MAP[id]
|
|
}
|
|
|
|
export type Model = BedrockModel | OpenAIModelId
|
|
|
|
export type ProviderModelConfig = {
|
|
/** Optional providerOptions to attach to the system message for this model */
|
|
systemProviderOptions?: Record<string, any>
|
|
/** The default model for this provider (used when limited or no preferred specified) */
|
|
default: boolean
|
|
}
|
|
|
|
export type ProviderRegistry = {
|
|
bedrock: {
|
|
models: Record<BedrockModel, ProviderModelConfig>
|
|
providerOptions?: Record<string, any>
|
|
}
|
|
openai: {
|
|
models: Record<OpenAIModelId, ProviderModelConfig>
|
|
providerOptions?: Record<string, any>
|
|
}
|
|
}
|
|
|
|
export const PROVIDERS: ProviderRegistry = {
|
|
bedrock: {
|
|
models: {
|
|
'anthropic.claude-3-7-sonnet-20250219-v1:0': {
|
|
systemProviderOptions: {
|
|
bedrock: {
|
|
// Always cache the system prompt (must not contain dynamic content)
|
|
cachePoint: { type: 'default' },
|
|
},
|
|
},
|
|
default: false,
|
|
},
|
|
'openai.gpt-oss-120b-1:0': {
|
|
default: true,
|
|
},
|
|
},
|
|
},
|
|
openai: {
|
|
models: {
|
|
'gpt-5.6-luna': { default: false },
|
|
'gpt-5.3-codex': { default: false },
|
|
'gpt-5.4-nano': { default: true },
|
|
},
|
|
providerOptions: {
|
|
openai: {
|
|
store: false,
|
|
},
|
|
},
|
|
},
|
|
}
|
|
|
|
export function getDefaultModelForProvider(provider: ProviderName): Model | undefined {
|
|
const models = PROVIDERS[provider]?.models as Record<Model, ProviderModelConfig>
|
|
if (!models) return undefined
|
|
|
|
return Object.keys(models).find((id) => models[id as Model]?.default) as Model | undefined
|
|
}
|