Files
OmniRoute/open-sse/services/ccOpenAiMediaBlocks.ts
Diego Rodrigues de Sa e Souza b2efa35982 fix(sse): CC bridge loses OpenAI-format image input (OpenCode/Kilo/Cline → AgentRouter) (#7888)
* fix(sse): convert OpenAI media parts to Claude blocks in the CC bridge

OpenAI-format clients (OpenCode/Kilo/Cline) reach the Claude-Code-compatible
bridge untranslated: chatCore skips the OpenAI->Claude translator when
sourceFormat is OPENAI, so image_url / AI-SDK image / file parts either went
upstream in OpenAI shape (silently ignored) or were dropped by the text-only
extraction, and media-only user turns were removed by hasValidContent().

- claudeCodeCompatible: convertOpenAiMediaBlock() converts image_url
  (base64 + remote), AI-SDK string image and file parts (pdf->document,
  image mime->image) to Claude blocks in both bridge paths; Claude-native
  blocks pass through unchanged and the text-only wire image is preserved.
- claudeHelper: hasValidContent() now counts image/document blocks so
  media-only user turns are not silently deleted.

Reported-by: beingshafin
Refs #7777

* refactor(sse): extract CC media-block conversion to ccOpenAiMediaBlocks.ts

claudeCodeCompatible.ts is frozen at 1202 lines by check:file-size; the #7777
helpers pushed it to 1291. Move them to a dedicated module, no behavior change.

* test(quality): register cc-bridge-openai-image-7777 test in stryker tap.testFiles
2026-07-20 15:56:40 -03:00

109 lines
3.6 KiB
TypeScript

/**
* OpenAI-format media parts reach the Claude-Code-compatible bridge
* untranslated when the source client speaks OpenAI (chatCore skips the
* OpenAI→Claude translator for CC-compatible providers), so `image_url` /
* AI SDK `image` / Chat Completions `file` parts must become Claude blocks
* before dispatch or the upstream silently ignores them (#7777).
* Claude-native blocks (`image`/`document` carrying `source`) are left for
* the caller to pass through unchanged.
*/
const DATA_URL_BASE64_PATTERN = /^data:([^;]+);base64,(.+)$/;
/** Returns a Claude block for an OpenAI media part, or null for non-media blocks. */
export function convertOpenAiMediaBlock(
record: Record<string, unknown>
): Record<string, unknown> | null {
if (record.type === "image_url") {
const rawUrl =
typeof record.image_url === "string" ? record.image_url : readRecord(record.image_url)?.url;
return claudeImageBlockFromUrl(rawUrl);
}
if (record.type === "image" && !readRecord(record.source) && typeof record.image === "string") {
return claudeImageBlockFromUrl(record.image);
}
if (record.type === "file") {
const file = readRecord(record.file);
const fileData = toNonEmptyString(file?.file_data) || toNonEmptyString(file?.data);
if (!fileData) return null;
const title = toNonEmptyString(file?.filename);
const match = fileData.match(DATA_URL_BASE64_PATTERN);
if (match) {
const mediaType = match[1];
const source = { type: "base64", media_type: mediaType, data: match[2] };
if (mediaType === "application/pdf") {
return { type: "document", source, ...(title ? { title } : {}) };
}
if (mediaType.startsWith("image/")) {
return { type: "image", source };
}
return null;
}
if (/^https?:\/\//i.test(fileData)) {
return {
type: "document",
source: { type: "url", url: fileData },
...(title ? { title } : {}),
};
}
return null;
}
return null;
}
/**
* Collects the Claude-shaped media blocks of a message's content array:
* OpenAI media parts are converted, Claude-native `image`/`document` blocks
* are cloned through as-is, everything else is ignored.
*/
export function collectClaudeMediaBlocks(content: unknown): Array<Record<string, unknown>> {
if (!Array.isArray(content)) return [];
const blocks: Array<Record<string, unknown>> = [];
for (const part of content) {
const record = readRecord(cloneValue(part));
if (!record) continue;
const media = convertOpenAiMediaBlock(record);
if (media) {
blocks.push(media);
continue;
}
if ((record.type === "image" || record.type === "document") && readRecord(record.source)) {
blocks.push(record);
}
}
return blocks;
}
function claudeImageBlockFromUrl(rawUrl: unknown): Record<string, unknown> | null {
const url = toNonEmptyString(rawUrl);
if (!url) return null;
const match = url.match(DATA_URL_BASE64_PATTERN);
if (match) {
return { type: "image", source: { type: "base64", media_type: match[1], data: match[2] } };
}
return { type: "image", source: { type: "url", url } };
}
function readRecord(value: unknown): Record<string, unknown> | null {
return value && typeof value === "object" && !Array.isArray(value)
? (value as Record<string, unknown>)
: null;
}
function toNonEmptyString(value: unknown): string | null {
if (typeof value !== "string") return null;
const trimmed = value.trim();
return trimmed || null;
}
function cloneValue<T>(value: T): T {
if (typeof structuredClone === "function") {
return structuredClone(value);
}
return JSON.parse(JSON.stringify(value)) as T;
}