Resolve conflicts: - streamingHandler.js: adopt upstreamResponseHeaders while keeping 0-token detail row avoidance - capabilities.js: preserve user-asserted caps and globalThis slots without local caching of catalogSource - AddCustomModelModal.js & providers/[id]/page.js: wire STT transport marker with custom model edits/assertions - models/custom/route.js & aliasRepo.js: persist custom model transport and invalidate user caps - usageRepo.js: key byApiKey live stats by full API key and keep tail in maskApiKey - UsageStats.js: lazy load charts dynamically
90 lines
5.2 KiB
JavaScript
90 lines
5.2 KiB
JavaScript
// Resolve valid thinking levels per model — drives UI level picker (suffix "model(level)").
|
|
// Reuses capabilities.js (thinkingFormat/canDisable) so this file only maps format→levels (DRY).
|
|
import { getCapabilitiesForModel } from "./capabilities.js";
|
|
import { matchPattern } from "./pricing.js";
|
|
import { resolveKiroEffortPath } from "../config/kiroConstants.js";
|
|
import { getProviderModels } from "../config/providerModels.js";
|
|
|
|
// Shared level sets (deduped) — verified against provider docs + wire in thinkingUnified.applyFormat.
|
|
const L = {
|
|
base: ["none", "low", "medium", "high"], // qwen, step, hunyuan, gemini-budget
|
|
onOff: ["none", "thinking"], // zai (binary), minimax (adaptive)
|
|
openai: ["none", "minimal", "low", "medium", "high", "xhigh"], // GPT-5.x / o-series (no "max")
|
|
levelMax: ["none", "low", "medium", "high", "max"], // claude-adaptive, kimi
|
|
budgetX: ["none", "low", "medium", "high", "xhigh", "max"], // claude-budget
|
|
gemini: ["minimal", "low", "medium", "high"], // gemini-3 thinkingLevel (no disable)
|
|
hiMax: ["none", "high", "max"], // deepseek (low/med→high, xhigh→max)
|
|
};
|
|
|
|
// thinkingFormat → valid selectable levels (source of truth for UI options).
|
|
const FORMAT_LEVELS = {
|
|
openai: L.openai,
|
|
"claude-adaptive": L.levelMax,
|
|
"claude-budget": L.budgetX,
|
|
"gemini-level": L.gemini,
|
|
"gemini-budget": L.base,
|
|
zai: L.onOff,
|
|
qwen: L.base,
|
|
kimi: L.levelMax,
|
|
deepseek: L.hiMax,
|
|
commandcode: ["none", "low", "medium", "high", "xhigh", "max"],
|
|
minimax: L.onOff,
|
|
hunyuan: L.base,
|
|
step: L.base,
|
|
};
|
|
|
|
const CODEX_GPT_5_6_LEVELS = ["none", "minimal", "low", "medium", "high", "xhigh", "max"];
|
|
|
|
// Model-name pattern overrides (glob, first match wins) — more precise than format default.
|
|
const PATTERN_THINKING = [
|
|
{ provider: "codex", pattern: "*gpt-6*", levels: CODEX_GPT_5_6_LEVELS },
|
|
{ provider: "codex", pattern: "*gpt-5.6-sol*", levels: [...CODEX_GPT_5_6_LEVELS, "ultra"] },
|
|
{ provider: "codex", pattern: "*gpt-5.6-terra*", levels: [...CODEX_GPT_5_6_LEVELS, "ultra"] },
|
|
{ provider: "codex", pattern: "*gpt-5.6-luna*", levels: CODEX_GPT_5_6_LEVELS },
|
|
{ pattern: "*codex*", levels: ["low", "medium", "high", "xhigh"] }, // codex cannot disable thinking
|
|
{ pattern: "*mimo*v2.6*", levels: ["none", "low", "medium", "high", "xhigh"] },
|
|
// mimo-v2.5-pro on opencode-go rejects reasoning_effort "max" (probed live); v2.5 accepts it.
|
|
{ pattern: "*mimo*v2.5-pro*", levels: ["none", "low", "medium", "high", "xhigh"] },
|
|
// DeepSeek v4.* (Alibaba MaaS, probed live): effort low|medium|high|xhigh|max
|
|
// all 200 via output_config.effort; "none" is a 400 on the anthropic route
|
|
// (disable thinking instead). none kept for the picker = disable.
|
|
{ pattern: "*deepseek-v4.*", levels: ["none", "low", "medium", "high", "xhigh", "max"] },
|
|
// codebuddy-cn per-model effort sets — the server's product-config payload
|
|
// publishes `reasoning.supportedEfforts` per model. NOTE: the chat endpoint
|
|
// accepts any level you send (probed none/minimal/low/medium/high/xhigh/max
|
|
// → all 200), but values outside a model's supportedEfforts are silently
|
|
// clamped, so the declared set stays authoritative for the picker. Models
|
|
// that publish no supportedEfforts (glm-5.1 / glm-5v-turbo / kimi-k2.x /
|
|
// kimi-k3-1 / minimax-m3) fall through to the openai format default.
|
|
{ provider: "codebuddy-cn", pattern: "glm-5.3*", levels: ["low", "high", "max"] },
|
|
{ provider: "codebuddy-cn", pattern: "glm-5.2", levels: ["high", "xhigh"] },
|
|
{ provider: "codebuddy-cn", pattern: "deepseek-v4*", levels: ["low", "high", "xhigh"] },
|
|
{ provider: "codebuddy-cn", pattern: "hy3*", levels: ["low", "high"] },
|
|
{ provider: "codebuddy-cn", pattern: "hy4*", levels: ["high"] },
|
|
// codebuddy-intl rides the same gateway catalog, so its deepseek levels match.
|
|
{ provider: "codebuddy-intl", pattern: "deepseek-v4*", levels: ["low", "high", "xhigh"] },
|
|
];
|
|
|
|
// The generic level set used when a model's thinking format is unknown. Exported
|
|
// for UI callers that must list levels for a model whose reasoning capability is
|
|
// user-asserted: the browser bundle has no access to the capability store, so it
|
|
// cannot derive a format the way getThinkingLevels does on the server.
|
|
export const BASE_THINKING_LEVELS = L.base;
|
|
|
|
// Returns valid thinking levels for a model, or null when the model has no reasoning.
|
|
export function getThinkingLevels(provider, model) {
|
|
if (provider === "kiro" && resolveKiroEffortPath(model) === null) return null;
|
|
const caps = getCapabilitiesForModel(provider, model);
|
|
if (!caps.reasoning) return null;
|
|
const baseId = String(model || "").replace(/\([^()]+\)\s*$/, "");
|
|
const modelLevels = provider === "codex"
|
|
? getProviderModels("cx").find((entry) => entry.id === baseId)?.thinkingLevels
|
|
: null;
|
|
const hit = PATTERN_THINKING.find((entry) =>
|
|
(!entry.provider || entry.provider === provider) && matchPattern(entry.pattern, model)
|
|
);
|
|
let levels = modelLevels || hit?.levels || FORMAT_LEVELS[caps.thinkingFormat] || L.base;
|
|
if (caps.thinkingCanDisable === false) levels = levels.filter((l) => l !== "none");
|
|
return levels;
|
|
}
|