Files
oh-my-pi/packages/coding-agent/src/utils/title-generator.ts
T
can1357 ef7636805b feat(coding-agent): removed canonical model variant selection and tracking
- Removed the canonical model variant indexing, selection, and tracking logic from the model registry and resolver.
- Eliminated the `canonical` sub-command, tab view, search tokens, and equivalence configuration structures from the CLI and model selector components.
- Refined model identification, lookup, and provider fallback resolution to bind exclusively to standard, raw model IDs.
- Relocated the equivalence utility script within the catalog package to support script-only policy generation.
2026-07-01 05:22:42 +02:00

349 lines
13 KiB
TypeScript

/**
* Generate session titles using a smol, fast model.
*/
import * as path from "node:path";
import { type Api, type AssistantMessage, completeSimple, type Model, type Tool } from "@oh-my-pi/pi-ai";
import { isTerminalHeadless, logger, prompt } from "@oh-my-pi/pi-utils";
import type { ModelRegistry } from "../config/model-registry";
import { resolveRoleSelection } from "../config/model-resolver";
import type { Settings } from "../config/settings";
import titleMarkerInstruction from "../prompts/system/title-marker-instruction.md" with { type: "text" };
import titleSystemPrompt from "../prompts/system/title-system.md" with { type: "text" };
import titleMarkerSystemPrompt from "../prompts/system/title-system-marker.md" with { type: "text" };
import { isTinyTitleLocalModelKey, ONLINE_TINY_TITLE_MODEL_KEY } from "../tiny/models";
import { formatTitleUserMessage, isLowSignalTitleInput, normalizeGeneratedTitle } from "../tiny/text";
import { tinyTitleClient } from "../tiny/title-client";
const TITLE_SYSTEM_PROMPT = prompt.render(titleSystemPrompt);
const TITLE_MARKER_SYSTEM_PROMPT = prompt.render(titleMarkerSystemPrompt);
const TITLE_MARKER_INSTRUCTION = prompt.render(titleMarkerInstruction);
const DEFAULT_TERMINAL_TITLE = "π";
const TERMINAL_TITLE_CONTROL_CHARS = /[\u0000-\u001f\u007f-\u009f]/g;
const TITLE_MAX_TOKENS = 30;
const REASONING_SAFE_MAX_TOKENS = 1024;
const SET_TITLE_TOOL_NAME = "set_title";
const setTitleTool: Tool = {
name: SET_TITLE_TOOL_NAME,
description: "Set the generated session title.",
parameters: {
type: "object",
properties: {
title: {
type: "string",
description:
'The generated session title, or exactly "none" when the message carries no concrete task yet.',
},
},
required: ["title"],
additionalProperties: false,
},
};
/** Matches the title a tool-choice-less model wraps in `<title>...</title>`. */
const TITLE_MARKER_RE = /<title>([\s\S]*?)<\/title>/i;
/**
* Whether the model honors a forced `tool_choice` so the `set_title` tool can be
* required. Providers/models that reject forced tool calls (chat-completions
* hosts without `tool_choice` support, Claude Fable/Mythos) can't be made to
* emit a structured call, so the caller falls back to marker-wrapped text.
*/
function modelSupportsForcedToolChoice(model: Model<Api>): boolean {
// `compat` is a union across APIs and `supportsToolChoice` lives only on the
// OpenAI-completions variant, so read both flags through a structural view.
const compat = model.compat as { supportsToolChoice?: boolean; supportsForcedToolChoice?: boolean } | undefined;
if (!compat) return true;
// A forced tool call first requires sending `tool_choice` at all. Hosts that
// drop the parameter entirely (`supportsToolChoice: false`, e.g. direct
// DeepSeek reasoning) can never be forced even when they otherwise accept
// forced values, so this veto wins over `supportsForcedToolChoice`.
if (compat.supportsToolChoice === false) return false;
if (typeof compat.supportsForcedToolChoice === "boolean") return compat.supportsForcedToolChoice;
if (typeof compat.supportsToolChoice === "boolean") return compat.supportsToolChoice;
return true;
}
function getTitleModel(registry: ModelRegistry, settings: Settings, currentModel?: Model<Api>): Model<Api> | undefined {
const availableModels = registry.getAvailable();
if (availableModels.length === 0) return undefined;
const titleModel = resolveRoleSelection(["tiny", "commit", "smol"], settings, availableModels)?.model;
if (titleModel) return titleModel;
if (currentModel) return currentModel;
return undefined;
}
/**
* Generate a title for a session based on the first user message.
*
* @param firstMessage The first user message
* @param registry Model registry
* @param settings Settings used to resolve the smol role
* @param sessionId Optional session id for sticky API key selection
* @param currentModel Current model (used to derive title model)
* @param metadataResolver Optional resolver evaluated after credential selection
* to produce request metadata (e.g. user_id for session attribution). Using a
* resolver instead of a pre-evaluated value ensures the metadata's account_uuid
* reflects the credential actually selected for this request.
* @param customSystemPrompt Optional title-specific system prompt override
*/
export async function generateSessionTitle(
firstMessage: string,
registry: ModelRegistry,
settings: Settings,
sessionId?: string,
currentModel?: Model<Api>,
metadataResolver?: (provider: string) => Record<string, unknown> | undefined,
customSystemPrompt?: string,
): Promise<string | null> {
// Defer titling for greetings / acknowledgements / empty input. The default
// tiny title model can't reliably decline trivial input, so this happens
// deterministically before any model is invoked; the caller retries on the
// next user message while the session stays unnamed.
if (isLowSignalTitleInput(firstMessage)) {
logger.debug("title-generator: skipped low-signal input", { sessionId, reason: "low-signal" });
return null;
}
const titleSystemPrompt = customSystemPrompt?.trim() || undefined;
const tinyModel = settings.get("providers.tinyModel");
if (tinyModel === ONLINE_TINY_TITLE_MODEL_KEY) {
return generateTitleOnline(
firstMessage,
registry,
settings,
sessionId,
currentModel,
metadataResolver,
undefined,
titleSystemPrompt,
);
}
// User explicitly picked a local tiny model. NEVER fall back to the online
// smol path (issue #3187): the smol role resolves through priority.json and
// silently bills whatever provider holds the resolved API key — OpenRouter
// in the reporter's case, leaking real credits without consent. If the
// local worker fails (unknown key, download missing, transformers.js
// crash, abort), leave the session untitled; the next user turn retries.
if (!isTinyTitleLocalModelKey(tinyModel)) {
logger.warn("title-generator: unknown local tiny model; skipping title (will not fall back to online)", {
sessionId,
model: tinyModel,
reason: "unknown-local-model",
});
return null;
}
try {
const localTitle = titleSystemPrompt
? await tinyTitleClient.generate(tinyModel, firstMessage, { systemPrompt: titleSystemPrompt })
: await tinyTitleClient.generate(tinyModel, firstMessage);
if (!localTitle) {
logger.warn("title-generator: local tiny model produced no title; skipping (no online fallback)", {
sessionId,
model: tinyModel,
reason: "local-no-output",
});
return null;
}
return localTitle;
} catch (err) {
logger.warn("title-generator: local tiny model errored; skipping (no online fallback)", {
sessionId,
model: tinyModel,
error: err instanceof Error ? err.message : String(err),
});
return null;
}
}
export async function generateTitleOnline(
firstMessage: string,
registry: ModelRegistry,
settings: Settings,
sessionId?: string,
currentModel?: Model<Api>,
metadataResolver?: (provider: string) => Record<string, unknown> | undefined,
signal?: AbortSignal,
customSystemPrompt?: string,
): Promise<string | null> {
const model = getTitleModel(registry, settings, currentModel);
if (!model) {
logger.warn("title-generator: no title model found", { sessionId, reason: "no-title-model" });
return null;
}
const titleSystemPrompt = customSystemPrompt?.trim() || undefined;
// Some providers can't be forced to call a tool — chat-completions hosts
// without `tool_choice` support, Claude Fable/Mythos — so a required
// `set_title` call never arrives. For those, ask the model to wrap the title
// in `<title>...</title>` markers and parse it from text instead.
const useForcedTool = modelSupportsForcedToolChoice(model);
const systemPrompt = useForcedTool
? [titleSystemPrompt ?? TITLE_SYSTEM_PROMPT]
: titleSystemPrompt
? [titleSystemPrompt, TITLE_MARKER_INSTRUCTION]
: [TITLE_MARKER_SYSTEM_PROMPT];
const userMessage = formatTitleUserMessage(firstMessage);
const modelName = `${model.provider}/${model.id}`;
const modelContext = {
sessionId,
provider: model.provider,
id: model.id,
model: modelName,
};
logger.debug("title-generator: start", modelContext);
try {
const apiKey = await registry.getApiKey(model, sessionId);
if (!apiKey) {
logger.warn("title-generator: no API key", { ...modelContext, reason: "missing-api-key" });
return null;
}
// Resolve metadata after getApiKey so the session-sticky credential for this
// request is already recorded; metadataResolver can then return the correct
// account_uuid rather than the snapshot-at-call-site value.
const metadata = metadataResolver?.(model.provider);
// Title generation is a 3-7 word task, but some reasoning backends ignore
// disableReasoning. Keep the normal cheap budget for non-reasoning models
// while reserving enough output room for reasoning models to still emit
// the forced tool call after any unavoidable thinking tokens.
const maxTokens = model.reasoning ? Math.max(TITLE_MAX_TOKENS, REASONING_SAFE_MAX_TOKENS) : TITLE_MAX_TOKENS;
logger.debug("title-generator: request", { ...modelContext, maxTokens });
const response = await completeSimple(
model,
{
systemPrompt,
messages: [{ role: "user", content: userMessage, timestamp: Date.now() }],
tools: useForcedTool ? [setTitleTool] : undefined,
},
{
apiKey: registry.resolver(model, sessionId),
maxTokens,
disableReasoning: true,
toolChoice: useForcedTool ? { type: "tool", name: SET_TITLE_TOOL_NAME } : undefined,
metadata,
signal,
},
);
if (response.stopReason === "error") {
logger.warn("title-generator: response error", {
...modelContext,
reason: "provider-response-error",
stopReason: response.stopReason,
errorMessage: response.errorMessage,
});
return null;
}
const title = normalizeGeneratedTitle(extractGeneratedTitle(response.content), firstMessage);
if (!title) {
logger.debug("title-generator: no title returned", {
...modelContext,
reason: "model-returned-none",
usage: response.usage,
stopReason: response.stopReason,
});
return null;
}
logger.debug("title-generator: success", {
...modelContext,
title,
usage: response.usage,
stopReason: response.stopReason,
});
return title;
} catch (err) {
logger.warn("title-generator: error", {
...modelContext,
reason: "exception",
error: err instanceof Error ? err.message : String(err),
});
return null;
}
}
function extractGeneratedTitle(contentBlocks: AssistantMessage["content"]): string {
let textTitle = "";
for (const content of contentBlocks) {
if (content.type === "toolCall" && content.name === SET_TITLE_TOOL_NAME) {
const args = content.arguments as Record<string, unknown>;
const title = args.title;
return typeof title === "string" ? title.trim() : "";
}
if (content.type === "text") {
textTitle += content.text;
}
}
// Tool-choice-less models are asked to wrap the title in <title>...</title>,
// but stay lenient: prefer the marker when the model closed it, otherwise
// accept a plain sentence after stripping any stray/unclosed tag fragment
// (e.g. output truncated before the closing tag).
const marker = TITLE_MARKER_RE.exec(textTitle);
if (marker) return marker[1].trim();
return textTitle.replace(/<\/?title>/gi, "").trim();
}
/**
* Remove control characters so model-generated titles cannot inject terminal escapes.
*/
function sanitizeTerminalTitlePart(value: string | undefined): string | undefined {
if (!value) return undefined;
const sanitized = value.replace(TERMINAL_TITLE_CONTROL_CHARS, "").trim();
return sanitized || undefined;
}
function getFallbackTerminalTitle(cwd: string | undefined): string | undefined {
if (!cwd) return undefined;
const resolvedCwd = path.resolve(cwd);
const baseName = path.basename(resolvedCwd);
if (!baseName || baseName === path.parse(resolvedCwd).root) return undefined;
return sanitizeTerminalTitlePart(baseName);
}
export function formatSessionTerminalTitle(sessionName: string | undefined, cwd?: string): string {
const label = sanitizeTerminalTitlePart(sessionName) ?? getFallbackTerminalTitle(cwd);
return label ? `${DEFAULT_TERMINAL_TITLE}: ${label}` : DEFAULT_TERMINAL_TITLE;
}
/**
* Set the terminal title using OSC 0 (sets both tab and window title). Unsupported terminals ignore it.
*/
export function setTerminalTitle(title: string): void {
if (!process.stdout.isTTY || isTerminalHeadless()) return;
process.stdout.write(`\x1b]0;${sanitizeTerminalTitlePart(title) ?? DEFAULT_TERMINAL_TITLE}\x07`);
}
export function setSessionTerminalTitle(sessionName: string | undefined, cwd?: string): void {
setTerminalTitle(formatSessionTerminalTitle(sessionName, cwd));
}
/**
* Save the current terminal title on terminals that support xterm window ops.
*/
export function pushTerminalTitle(): void {
if (!process.stdout.isTTY || isTerminalHeadless()) return;
process.stdout.write("\x1b[22;2t");
}
/**
* Restore the previously saved terminal title on terminals that support xterm window ops.
*/
export function popTerminalTitle(): void {
if (!process.stdout.isTTY || isTerminalHeadless()) return;
process.stdout.write("\x1b[23;2t");
}