Files
9router/tests/unit/opencode-muse-spark-thinking.test.js
anojndr ab044e6d6d fix(opencode): route Muse Spark through the Responses API
muse-spark-1.2-contributor-free returned HTTP 500 on /zen/v1/chat/completions.
The model is only served by /zen/v1/responses, so route it there via a per-model
targetFormat and normalize the Chat fields the Responses API rejects
(max_tokens -> max_output_tokens, reasoning_effort -> reasoning{effort,summary}),
clamping max/ultra down to the highest effort the model accepts (xhigh).

Routing stays per-model: the other free models (big-pickle, hy3-free, mimo,
nemotron, laguna) are not served by /responses and keep /chat/completions.
2026-08-28 11:33:14 +07:00

95 lines
3.1 KiB
JavaScript

import { describe, expect, it } from "vitest";
import { getCapabilitiesForModel } from "../../open-sse/providers/capabilities.js";
import { PROVIDER_MODELS } from "../../open-sse/config/providerModels.js";
import { getThinkingLevels } from "../../open-sse/providers/thinkingLevels.js";
import { FORMATS } from "../../open-sse/translator/formats.js";
import { OpenCodeExecutor } from "../../open-sse/executors/opencode.js";
import "../translator/registerAll.js";
import { translateRequest } from "../../open-sse/translator/index.js";
const MODEL = "muse-spark-1.2-contributor-free";
const PROVIDER = "opencode";
const input = [{
type: "message",
role: "user",
content: [{ type: "input_text", text: "Think, then answer: 2 + 2?" }],
}];
describe("OpenCode Free Muse Spark thinking", () => {
it("advertises reasoning and the requested model limits", () => {
expect(PROVIDER_MODELS.oc?.some((model) => model.id === MODEL)).toBe(true);
expect(getCapabilitiesForModel(PROVIDER, MODEL)).toMatchObject({
reasoning: true,
thinkingFormat: "openai",
contextWindow: 1048576,
maxOutput: 131072,
});
expect(getCapabilitiesForModel(PROVIDER, `oc/${MODEL}`)).toMatchObject({
reasoning: true,
contextWindow: 1048576,
maxOutput: 131072,
});
expect(getThinkingLevels(PROVIDER, MODEL)).toEqual([
"none",
"minimal",
"low",
"medium",
"high",
"xhigh",
]);
});
it("clamps max to xhigh and emits the Responses reasoning shape", () => {
const body = {
input,
reasoning: { effort: "max" },
max_tokens: 131072,
};
const out = new OpenCodeExecutor().transformRequest(MODEL, body, true, {
connectionId: "opencode-muse-spark-test",
});
expect(out.reasoning).toEqual({ effort: "xhigh", summary: "auto" });
expect(out.reasoning_effort).toBeUndefined();
expect(out.max_output_tokens).toBe(131072);
expect(out.max_tokens).toBeUndefined();
});
it("leaves the other free models on Chat Completions", () => {
const executor = new OpenCodeExecutor();
const body = { messages: [{ role: "user", content: "hi" }], max_tokens: 1024 };
executor.transformRequest("big-pickle", body, true, {});
expect(executor.buildUrl("big-pickle")).toBe("https://opencode.ai/zen/v1/chat/completions");
expect(body.max_tokens).toBe(1024);
expect(body.max_output_tokens).toBeUndefined();
});
it("translates Chat Completions max thinking into a Responses request", () => {
const body = {
model: `oc/${MODEL}`,
messages: [{ role: "user", content: "Think, then answer: 2 + 2?" }],
reasoning_effort: "max",
max_tokens: 131072,
};
const translated = translateRequest(
FORMATS.OPENAI,
FORMATS.OPENAI_RESPONSES,
MODEL,
body,
true,
{},
PROVIDER,
);
const out = new OpenCodeExecutor().transformRequest(MODEL, translated, true, {
connectionId: "opencode-muse-spark-translation-test",
});
expect(out.reasoning).toEqual({ effort: "xhigh", summary: "auto" });
expect(out.max_output_tokens).toBe(131072);
expect(out.max_tokens).toBeUndefined();
});
});