fix(catalog): kept qwen3.8 max preview on enable_thinking

Restored the preview to its main compat so the reasoning_effort dialect is scoped to qwen3.8-max, leaving the preview ladder unchanged.
This commit is contained in:
roboomp
2026-08-08 14:54:29 +00:00
parent 4d2c6e37f1
commit c0eda613c8
3 changed files with 19 additions and 3 deletions
@@ -175,4 +175,18 @@ describe("OpenAI compat policy", () => {
expect(params.chat_template_kwargs).toBeUndefined();
}
});
it("keeps Token Plan qwen3.8-max-preview on the enable_thinking dialect", () => {
// The preview rides Alibaba's binary enable_thinking toggle, not the
// OpenAI reasoning_effort control, so effort selections must not leak an
// unsupported reasoning_effort onto the wire.
const model = getBundledModel<"openai-completions">("alibaba-token-plan", "qwen3.8-max-preview");
const params = chatParams();
applyChatCompletionsCompatPolicy(
params,
resolveOpenAICompatPolicy(model, { endpoint: "chat-completions", reasoning: Effort.High }),
);
expect(params.enable_thinking).toBe(true);
expect(params.reasoning_effort).toBeUndefined();
});
});
+1 -2
View File
@@ -8070,8 +8070,7 @@
},
"compat": {
"supportsDeveloperRole": false,
"supportsReasoningEffort": true,
"thinkingFormat": "openai"
"supportsReasoningEffort": true
}
}
},
@@ -2714,7 +2714,10 @@ export const ALIBABA_TOKEN_PLAN_STATIC_MODELS: readonly ModelSpec<"openai-comple
efforts: [Effort.Low, Effort.High, Effort.XHigh],
requiresEffort: true,
},
compat: ALIBABA_TOKEN_PLAN_QWEN_EFFORT_COMPAT,
compat: {
...ALIBABA_TOKEN_PLAN_COMPAT,
supportsReasoningEffort: true,
},
},
{
id: "qwen3.8-max",