diff --git a/scripts/ab-corpus/clients/minimax.ts b/scripts/ab-corpus/clients/minimax.ts index d40587c4f..e6303e736 100644 --- a/scripts/ab-corpus/clients/minimax.ts +++ b/scripts/ab-corpus/clients/minimax.ts @@ -33,7 +33,16 @@ export async function callMinimax(args: CallMinimaxArgs): Promise blocks in output-parser.ts); the openai-compat + // default 4096 is fine for obvious prompts but cuts close on + // composite multi-tool outputs (12-tag responses + thinking can + // approach 3-4k easy). Double the cap defensively — we already pay + // for thinking either way, and most replies still come in well + // under 1k completion tokens (ab-v4 avg 697), so the bigger cap + // costs nothing on the happy path and only matters when the model + // would otherwise truncate mid-output. + maxTokens: args.maxTokens ?? 8192, label: 'minimax', // ab-v3 saw 4 minimax timeouts on 104 runs — all wall-clock aborts, // none were content-quality issues. One retry picks up the