fix(catalog): corrected qwen3.8 max discovery metadata

Curated reasoning, multimodal input, context limits, and the provider-specific effort ladder for the discovered Alibaba Token Plan model.

Fixes #8019
This commit is contained in:
roboomp
2026-08-08 14:37:24 +00:00
parent 08819b279c
commit 155fdaedba
7 changed files with 5385 additions and 2090 deletions
@@ -172,7 +172,13 @@ export function applyGeneratedModelPolicies(models: ModelSpec<Api>[]): void {
*/
export function rebakeModelThinking(model: ModelSpec<Api>): void {
if (isVariantCollapsedSpec(model)) return;
if (model.provider === "alibaba-token-plan" && model.id === "qwen3.8-max-preview" && model.thinking) return;
if (
model.provider === "alibaba-token-plan" &&
(model.id === "qwen3.8-max-preview" || model.id === "qwen3.8-max") &&
model.thinking
) {
return;
}
const requiresProviderAuthoredEffort =
model.provider === "umans" && (model.thinking?.requiresEffort === true || model.id === "umans-kimi-k2.7");
const thinking = resolveModelThinking({ ...model, thinking: undefined }, buildCompat(model));