mirror of
https://github.com/heygen-com/hyperframes.git
synced 2026-09-14 18:01:20 +08:00
a4303137cb
* fix(skills): storyboard review — bg-on-clip rule, slideshow output, parser parity guard Addresses the storyboard-angle review (jrusso1020): - M1 (invisible text): frame-worker.md (x3) + SKILL.md Step 5 (x3) now require a frame's full-bleed background on a class=clip layer, never the #root / data-composition-id element (the root is clip-gated to its scene window, so a background on it is not a dependable ground and dark text can land on the black host body). The assembler already paints frame.md's canvas onto index #root as the base ground; the per-frame clip rides on top. - B3 (slideshow truncates to slide 1): slideshow/SKILL.md gains an Output section (decks render via 'present'; 'render index.html' captures only the first composition; linear main-line MP4 export is deferred). - Parser drift: vendoredParity.test.ts guards the three vendored storyboard.mjs copies (byte-identical + parse-parity with @hyperframes/core). - skills-manifest.json regenerated for the edited SKILL.md files. Co-Authored-By: Claude Opus 4.8 (1M context) <noreply@anthropic.com> * fix(cli): storyboard review — lint, validate help, snapshot, inspect, capture, render Addresses the CLI findings from the storyboard-angle review (jrusso1020): - lint (@hyperframes/lint): accept vendor-prefixed system-font keywords -apple-system / BlinkMacSystemFont so a system stack with a generic fallback no longer trips font_family_without_font_face (+ test). - help: list 'validate' under Project in 'hyperframes --help' (was runnable but undocumented). - snapshot: honor -o/--output (the flag did not exist; output was hardcoded to snapshots/). The dir is resolved once and threaded through capture + contact sheet + Gemini. - snapshot: split font status into loaded / error / unused with a one-line summary; only a real 'error' is reported as FAILED (an unrequested @font-face is 'unused', not a contradiction with 'loaded'). - inspect: suppress text_occluded across a scene-to-scene crossfade (occluder in a different data-composition-id mount while a scene is mid-fade); a same-scene or two-settled-scenes overlap still flags. - inspect: suppress content_overlap between in-flow siblings governed by the same flex/grid container (tight stacks / number lockups are layout slop). - capture: record source resolution (videoWidth/Height) in video-manifest.json alongside the DOM display box; consumers size off the source dims. - render: warn when the target carries a slideshow island (render captures only the first scene, so the MP4 is truncated to slide 1; use 'present'). Co-Authored-By: Claude Opus 4.8 (1M context) <noreply@anthropic.com> --------- Co-authored-by: Claude Opus 4.8 (1M context) <noreply@anthropic.com>
534 lines
22 KiB
TypeScript
534 lines
22 KiB
TypeScript
/**
|
||
* Content extraction helpers for the website capture pipeline.
|
||
*
|
||
* Handles library detection, visible text extraction, vision captioning,
|
||
* and asset description generation.
|
||
*
|
||
* All page.evaluate() calls use string expressions to avoid
|
||
* tsx/esbuild __name injection (see esbuild issue #1031).
|
||
*/
|
||
|
||
import type { Page } from "puppeteer-core";
|
||
import { existsSync, readdirSync, statSync, readFileSync } from "node:fs";
|
||
import { basename, join } from "node:path";
|
||
import type sharpType from "sharp";
|
||
import type { CatalogedAsset } from "./assetCataloger.js";
|
||
import type { DesignTokens } from "./types.js";
|
||
|
||
/**
|
||
* Detect JS libraries via window globals, DOM fingerprints, script URLs,
|
||
* and WebGL shader analysis.
|
||
*
|
||
* Returns a deduplicated list of detected library names.
|
||
*/
|
||
export async function detectLibraries(
|
||
page: Page,
|
||
capturedShaders?: Array<{ type: string; source: string }>,
|
||
): Promise<string[]> {
|
||
let detectedLibraries: string[] = [];
|
||
try {
|
||
detectedLibraries = (await page.evaluate(`(() => {
|
||
var libs = [];
|
||
function add(name) { if (libs.indexOf(name) === -1) libs.push(name); }
|
||
|
||
// 1. Window globals (works for CDN-loaded / non-bundled libraries)
|
||
if (typeof window.gsap !== 'undefined' || typeof window.TweenMax !== 'undefined') add('GSAP');
|
||
if (typeof window.ScrollTrigger !== 'undefined') add('GSAP ScrollTrigger');
|
||
if (typeof window.THREE !== 'undefined') add('Three.js');
|
||
if (typeof window.PIXI !== 'undefined') add('PixiJS');
|
||
if (typeof window.BABYLON !== 'undefined') add('Babylon.js');
|
||
if (typeof window.Lottie !== 'undefined' || typeof window.lottie !== 'undefined') add('Lottie');
|
||
if (typeof window.__NEXT_DATA__ !== 'undefined') add('Next.js');
|
||
if (typeof window.__NUXT__ !== 'undefined') add('Nuxt');
|
||
if (typeof window.Webflow !== 'undefined') add('Webflow');
|
||
|
||
// 2. DOM fingerprints (survive bundling — most reliable for modern sites)
|
||
// Three.js sets data-engine on every canvas it creates
|
||
var threeCanvas = document.querySelector('canvas[data-engine*="three"]');
|
||
if (threeCanvas) add('Three.js (' + (threeCanvas.getAttribute('data-engine') || '') + ')');
|
||
// Babylon.js also sets data-engine
|
||
var babylonCanvas = document.querySelector('canvas[data-engine*="Babylon"]');
|
||
if (babylonCanvas) add('Babylon.js');
|
||
// Lottie web components
|
||
if (document.querySelector('dotlottie-wc, lottie-player, dotlottie-player')) add('Lottie');
|
||
// Rive
|
||
if (document.querySelector('canvas[class*="rive"], rive-canvas')) add('Rive');
|
||
// React/Next.js
|
||
if (document.getElementById('__next')) add('Next.js');
|
||
if (document.getElementById('__nuxt')) add('Nuxt');
|
||
if (document.querySelector('[data-reactroot], [data-react-helmet]')) add('React');
|
||
// Svelte
|
||
if (document.querySelector('[class*="svelte-"]')) add('Svelte');
|
||
// Tailwind (utility class detection)
|
||
if (document.querySelector('[class*="flex "], [class*="grid "], [class*="px-"], [class*="py-"]')) add('Tailwind CSS');
|
||
// Framer Motion
|
||
if (document.querySelector('[style*="--framer-"], [data-framer-component-type]')) add('Framer Motion');
|
||
|
||
// 3. Script URL patterns
|
||
document.querySelectorAll('script[src]').forEach(function(s) {
|
||
var src = s.src.toLowerCase();
|
||
if (src.includes('gsap') || src.includes('tweenmax') || src.includes('greensock')) add('GSAP');
|
||
if (src.includes('scrolltrigger')) add('GSAP ScrollTrigger');
|
||
if (src.includes('three.module') || src.includes('three.min')) add('Three.js');
|
||
if (src.includes('pixi')) add('PixiJS');
|
||
if (src.includes('lottie') || src.includes('bodymovin')) add('Lottie');
|
||
if (src.includes('framer-motion')) add('Framer Motion');
|
||
if (src.includes('anime.min') || src.includes('animejs')) add('Anime.js');
|
||
if (src.includes('matter.min') || src.includes('matter-js')) add('Matter.js');
|
||
if (src.includes('lenis')) add('Lenis (smooth scroll)');
|
||
});
|
||
|
||
return libs;
|
||
})()`)) as string[];
|
||
} catch {
|
||
// Non-blocking
|
||
}
|
||
|
||
// 4. Shader fingerprinting — infer WebGL framework from captured GLSL
|
||
try {
|
||
const shaders = capturedShaders || [];
|
||
if (shaders.length > 0) {
|
||
const allSource = shaders.map((s) => s.source).join("\n");
|
||
const add = (name: string) => {
|
||
if (!detectedLibraries.includes(name)) detectedLibraries.push(name);
|
||
};
|
||
add("WebGL");
|
||
// Three.js shader fingerprints (built-in uniforms that survive bundling)
|
||
if (allSource.includes("modelViewMatrix") && allSource.includes("projectionMatrix"))
|
||
add("Three.js (confirmed via shaders)");
|
||
// PixiJS shader fingerprints
|
||
else if (
|
||
allSource.includes("vTextureCoord") &&
|
||
allSource.includes("uSampler") &&
|
||
!allSource.includes("modelViewMatrix")
|
||
)
|
||
add("PixiJS (confirmed via shaders)");
|
||
// Babylon.js shader fingerprints
|
||
else if (allSource.includes("viewProjection") && allSource.includes("world"))
|
||
add("Babylon.js (confirmed via shaders)");
|
||
}
|
||
} catch {
|
||
/* non-blocking */
|
||
}
|
||
|
||
return detectedLibraries;
|
||
}
|
||
|
||
/**
|
||
* Extract all visible text from the page in DOM order using a TreeWalker.
|
||
* Truncates to ~30K chars to avoid blowing up downstream prompts.
|
||
*/
|
||
export async function extractVisibleText(page: Page): Promise<string> {
|
||
let visibleTextContent = "";
|
||
try {
|
||
visibleTextContent = (await page.evaluate(`(() => {
|
||
var cookieRe = /^(accept|cookie|privacy|that's fine|got it|i agree|reject all|accept all|manage cookies|consent)$/i;
|
||
var walker = document.createTreeWalker(document.body, NodeFilter.SHOW_TEXT, null);
|
||
var texts = [];
|
||
var node;
|
||
while (node = walker.nextNode()) {
|
||
var text = (node.textContent || '').trim();
|
||
if (text.length < 3) continue;
|
||
var el = node.parentElement;
|
||
if (!el) continue;
|
||
var style = getComputedStyle(el);
|
||
if (style.display === 'none' || style.visibility === 'hidden' || style.opacity === '0') continue;
|
||
var tag = el.tagName.toLowerCase();
|
||
if (tag === 'script' || tag === 'style' || tag === 'noscript') continue;
|
||
// Skip very short text inside nav/footer (catches single-word nav links)
|
||
// Threshold is 8 chars to preserve footer copy like "© 2026 Stripe" (16 chars)
|
||
var inNavOrFooter = el.closest('nav, footer, [role="navigation"]');
|
||
if (inNavOrFooter && text.length < 8) continue;
|
||
// Skip common cookie/consent patterns
|
||
if (cookieRe.test(text)) continue;
|
||
texts.push('[' + tag + '] ' + text);
|
||
}
|
||
return texts.join('\\n');
|
||
})()`)) as string;
|
||
// Truncate to ~30K chars to avoid blowing up the prompt
|
||
if (visibleTextContent.length > 30000) {
|
||
visibleTextContent = visibleTextContent.slice(0, 30000) + "\n[...truncated]";
|
||
}
|
||
} catch {
|
||
// Non-blocking
|
||
}
|
||
return visibleTextContent;
|
||
}
|
||
|
||
/**
|
||
* Caption downloaded images using a vision model.
|
||
*
|
||
* Provider is chosen by which API key is present: OPENROUTER_API_KEY → OpenRouter
|
||
* (any vision model via its OpenAI-style API), else GEMINI_API_KEY/GOOGLE_API_KEY
|
||
* → Google Gemini, else no captioning. OpenRouter wins if both are set.
|
||
*
|
||
* Batches requests to stay under free-tier rate limits.
|
||
* Returns a map of filename -> caption string.
|
||
*/
|
||
export async function captionImagesWithGemini(
|
||
outputDir: string,
|
||
progress: (stage: string, detail?: string) => void,
|
||
warnings: string[],
|
||
): Promise<Record<string, string>> {
|
||
const geminiCaptions: Record<string, string> = {};
|
||
const openRouterKey = process.env.OPENROUTER_API_KEY;
|
||
const geminiKey = process.env.GEMINI_API_KEY || process.env.GOOGLE_API_KEY;
|
||
if (!openRouterKey && !geminiKey) return geminiCaptions;
|
||
|
||
// OpenRouter takes priority when both keys are set — it's the explicit opt-in
|
||
// for users without Google access. Both providers satisfy the same
|
||
// single-image → one-line-caption contract (`captionOne`), so the batching and
|
||
// SVG-rasterization loops below stay provider-agnostic.
|
||
const useOpenRouter = Boolean(openRouterKey);
|
||
const providerName = useOpenRouter ? "OpenRouter" : "Gemini";
|
||
// Default mirrors the Gemini path's tier (3.x flash-lite). Override per
|
||
// provider via HYPERFRAMES_OPENROUTER_MODEL / HYPERFRAMES_GEMINI_MODEL.
|
||
const model = useOpenRouter
|
||
? process.env.HYPERFRAMES_OPENROUTER_MODEL || "google/gemini-3.1-flash-lite"
|
||
: process.env.HYPERFRAMES_GEMINI_MODEL || "gemini-3.1-flash-lite-preview";
|
||
|
||
progress("design", `Captioning images with ${providerName} vision...`);
|
||
try {
|
||
// One image → one short caption. Each provider implements this contract;
|
||
// everything below is provider-agnostic.
|
||
type CaptionOne = (args: {
|
||
mimeType: string;
|
||
base64: string;
|
||
prompt: string;
|
||
maxTokens: number;
|
||
}) => Promise<string>;
|
||
|
||
let captionOne: CaptionOne;
|
||
if (openRouterKey) {
|
||
captionOne = async ({ mimeType, base64, prompt, maxTokens }) => {
|
||
const res = await fetch("https://openrouter.ai/api/v1/chat/completions", {
|
||
method: "POST",
|
||
headers: {
|
||
Authorization: `Bearer ${openRouterKey}`,
|
||
"Content-Type": "application/json",
|
||
},
|
||
body: JSON.stringify({
|
||
model,
|
||
messages: [
|
||
{
|
||
role: "user",
|
||
content: [
|
||
{ type: "text", text: prompt },
|
||
{ type: "image_url", image_url: { url: `data:${mimeType};base64,${base64}` } },
|
||
],
|
||
},
|
||
],
|
||
max_tokens: maxTokens,
|
||
}),
|
||
});
|
||
if (!res.ok) {
|
||
const detail = await res.text().catch(() => "");
|
||
throw new Error(`OpenRouter ${res.status} ${res.statusText}: ${detail.slice(0, 200)}`);
|
||
}
|
||
const data = (await res.json()) as {
|
||
choices?: Array<{ message?: { content?: string } }>;
|
||
};
|
||
return data.choices?.[0]?.message?.content?.trim() || "";
|
||
};
|
||
} else {
|
||
// Unreachable when geminiKey is unset (guarded above); re-narrow for TS.
|
||
if (!geminiKey) return geminiCaptions;
|
||
const { GoogleGenAI } = await import("@google/genai");
|
||
const ai = new GoogleGenAI({ apiKey: geminiKey });
|
||
captionOne = async ({ mimeType, base64, prompt, maxTokens }) => {
|
||
const response = await ai.models.generateContent({
|
||
model,
|
||
contents: [
|
||
{ role: "user", parts: [{ inlineData: { mimeType, data: base64 } }, { text: prompt }] },
|
||
],
|
||
config: { maxOutputTokens: maxTokens },
|
||
});
|
||
return response.text?.trim() || "";
|
||
};
|
||
}
|
||
|
||
const imageFiles = readdirSync(join(outputDir, "assets")).filter((f: string) =>
|
||
/\.(png|jpg|jpeg|webp|gif)$/i.test(f),
|
||
);
|
||
|
||
// Caption in parallel batches. Gemini free tier is ~5 RPM (slow but $0),
|
||
// paid/OpenRouter ~2000 RPM. We batch 20 with a 2s inter-batch pause and rely
|
||
// on Promise.allSettled so a rate-limited image degrades to "" rather than
|
||
// failing the batch.
|
||
const BATCH_SIZE = 20;
|
||
for (let i = 0; i < imageFiles.length; i += BATCH_SIZE) {
|
||
const batch = imageFiles.slice(i, i + BATCH_SIZE);
|
||
const results = await Promise.allSettled(
|
||
batch.map(async (file: string) => {
|
||
const filePath = join(outputDir, "assets", file);
|
||
const stat = statSync(filePath);
|
||
if (stat.size > 4_000_000) return { file, caption: "" }; // skip images > 4 MB (provider inline limit)
|
||
const buffer = readFileSync(filePath);
|
||
const base64 = buffer.toString("base64");
|
||
const ext = file.split(".").pop()?.toLowerCase() || "png";
|
||
const mimeType = ext === "jpg" ? "image/jpeg" : `image/${ext}`;
|
||
const caption = await captionOne({
|
||
mimeType,
|
||
base64,
|
||
prompt:
|
||
"Describe this website image in ONE short sentence for a video storyboard. Focus on: what it shows, dominant colors, whether background is light or dark. Be factual, not creative.",
|
||
maxTokens: 500,
|
||
});
|
||
return { file, caption };
|
||
}),
|
||
);
|
||
for (const result of results) {
|
||
if (result.status === "fulfilled" && result.value.caption) {
|
||
geminiCaptions[result.value.file] = result.value.caption;
|
||
}
|
||
}
|
||
// Pace requests between batches (paid tier: 2000+ RPM, free tier: rate-limited)
|
||
if (i + BATCH_SIZE < imageFiles.length) {
|
||
await new Promise((r) => setTimeout(r, 2000)); // 2s pause between batches — paid tier handles 2000 RPM, free tier retries via Promise.allSettled
|
||
}
|
||
progress(
|
||
"design",
|
||
`Captioned ${Math.min(i + BATCH_SIZE, imageFiles.length)}/${imageFiles.length} images...`,
|
||
);
|
||
}
|
||
progress(
|
||
"design",
|
||
`${Object.keys(geminiCaptions).length} images captioned with ${providerName}`,
|
||
);
|
||
|
||
// Rasterize SVGs to PNG before captioning — Vision hallucinates wordmarks when reading SVG path text.
|
||
const svgFiles: Array<{ file: string; relPath: string }> = [];
|
||
const assetsDir = join(outputDir, "assets");
|
||
for (const f of readdirSync(assetsDir)) {
|
||
if (/\.svg$/i.test(f)) svgFiles.push({ file: f, relPath: f });
|
||
}
|
||
const svgsSubdir = join(assetsDir, "svgs");
|
||
if (existsSync(svgsSubdir)) {
|
||
for (const f of readdirSync(svgsSubdir)) {
|
||
if (/\.svg$/i.test(f)) svgFiles.push({ file: f, relPath: `svgs/${f}` });
|
||
}
|
||
}
|
||
|
||
if (svgFiles.length > 0) {
|
||
// sharp is an optional native module; its platform binary fails to load
|
||
// on some installs (omit-optional, musl/glibc, monorepo hoisting, broken
|
||
// cache). Load it lazily and degrade to skipping SVG captioning rather
|
||
// than crashing the whole capture command on import.
|
||
let sharp: typeof sharpType;
|
||
try {
|
||
sharp = (await import("sharp")).default as typeof sharpType;
|
||
} catch (err) {
|
||
warnings.push(
|
||
`Skipped ${svgFiles.length} SVG caption(s): sharp could not load (${(err as Error).message}). ` +
|
||
`Reinstall with optional dependencies enabled (e.g. \`npm i sharp\`) to caption SVG assets.`,
|
||
);
|
||
return geminiCaptions;
|
||
}
|
||
progress("design", `Rasterizing + captioning ${svgFiles.length} SVGs via vision API...`);
|
||
const SVG_BATCH = 20;
|
||
const SVG_RENDER_SIZE = 256; // px — enough resolution for Gemini to read wordmarks, small enough to keep payload sub-MB
|
||
let svgsSkipped = 0;
|
||
for (let i = 0; i < svgFiles.length; i += SVG_BATCH) {
|
||
const batch = svgFiles.slice(i, i + SVG_BATCH);
|
||
const results = await Promise.allSettled(
|
||
batch.map(async ({ relPath }) => {
|
||
const filePath = join(assetsDir, relPath);
|
||
let pngBase64: string;
|
||
try {
|
||
// Flatten against a contrasting background — white-on-white SVGs render invisible to Vision.
|
||
const svgSource = readFileSync(filePath, "utf-8");
|
||
const lightFillHits = (
|
||
svgSource.match(/fill\s*=\s*["'](#fff(fff)?|white|#[ef][ef][ef]|#[ef]{6})["']/gi) ||
|
||
[]
|
||
).length;
|
||
const darkFillHits = (
|
||
svgSource.match(/fill\s*=\s*["'](#000(000)?|black|#[0-3]{6}|#[0-3]{3})["']/gi) || []
|
||
).length;
|
||
const bg =
|
||
lightFillHits > darkFillHits
|
||
? { r: 32, g: 32, b: 32 } // dark slate behind light glyphs
|
||
: { r: 255, g: 255, b: 255 }; // white behind dark glyphs (default)
|
||
const pngBuffer = await sharp(filePath)
|
||
.resize({
|
||
width: SVG_RENDER_SIZE,
|
||
height: SVG_RENDER_SIZE,
|
||
fit: "inside",
|
||
withoutEnlargement: false,
|
||
})
|
||
.flatten({ background: bg })
|
||
.png()
|
||
.toBuffer();
|
||
pngBase64 = pngBuffer.toString("base64");
|
||
} catch {
|
||
// exotic SVG features may break sharp; skip caption rather than block
|
||
svgsSkipped++;
|
||
return { file: relPath, caption: "" };
|
||
}
|
||
const caption = await captionOne({
|
||
mimeType: "image/png",
|
||
base64: pngBase64,
|
||
prompt:
|
||
"Describe this SVG asset rendered from a website in ONE short sentence for a video storyboard. " +
|
||
"Focus on: what shape/icon/illustration/wordmark it is, its colors, any text it contains. " +
|
||
"If you see a wordmark, READ THE LETTERS LITERALLY — do not guess a brand from context. " +
|
||
"Be factual.",
|
||
maxTokens: 300,
|
||
});
|
||
return { file: relPath, caption };
|
||
}),
|
||
);
|
||
for (const result of results) {
|
||
if (result.status === "fulfilled" && result.value.caption) {
|
||
geminiCaptions[result.value.file] = result.value.caption;
|
||
}
|
||
}
|
||
if (i + SVG_BATCH < svgFiles.length) {
|
||
await new Promise((r) => setTimeout(r, 2000));
|
||
}
|
||
progress(
|
||
"design",
|
||
`Captioned ${Math.min(i + SVG_BATCH, svgFiles.length)}/${svgFiles.length} SVGs...`,
|
||
);
|
||
}
|
||
progress("design", `${Object.keys(geminiCaptions).length} total assets captioned`);
|
||
if (svgsSkipped > 0) {
|
||
progress(
|
||
"design",
|
||
`skipped rasterizing ${svgsSkipped} SVG(s) — fell back to label-derived`,
|
||
);
|
||
}
|
||
}
|
||
} catch (err) {
|
||
warnings.push(`${providerName} captioning failed: ${err}`);
|
||
}
|
||
|
||
return geminiCaptions;
|
||
}
|
||
|
||
/**
|
||
* Generate asset-descriptions.md — one-line descriptions for each downloaded asset.
|
||
*
|
||
* Returns the description lines (without the markdown header).
|
||
*/
|
||
export function generateAssetDescriptions(
|
||
outputDir: string,
|
||
tokens: DesignTokens,
|
||
catalogedAssets: CatalogedAsset[],
|
||
geminiCaptions: Record<string, string>,
|
||
): string[] {
|
||
// Sort: Gemini-captioned images first (richest descriptions), then uncaptioned, then SVGs, then fonts
|
||
const captionedLines: string[] = [];
|
||
const uncaptionedLines: string[] = [];
|
||
const svgLines: string[] = [];
|
||
const fontLines: string[] = [];
|
||
|
||
// Describe downloaded images
|
||
const assetsPath = join(outputDir, "assets");
|
||
try {
|
||
for (const file of readdirSync(assetsPath)) {
|
||
if (file === "svgs" || file === "fonts" || file === "lottie" || file === "videos") continue;
|
||
const filePath = join(assetsPath, file);
|
||
const stat = statSync(filePath);
|
||
if (!stat.isFile()) continue;
|
||
const sizeKb = Math.round(stat.size / 1024);
|
||
const catalogMatch = catalogedAssets.find(
|
||
(a) => a.url && file.includes(a.url.split("/").pop()?.split("?")[0]?.slice(0, 20) || "___"),
|
||
);
|
||
const desc = catalogMatch?.description || catalogMatch?.notes || "";
|
||
const heading = catalogMatch?.nearestHeading || "";
|
||
const section = catalogMatch?.sectionClasses || "";
|
||
const aboveFold = catalogMatch?.aboveFold ? "above fold" : "";
|
||
const geminiCaption = geminiCaptions[file];
|
||
const cleanName = file.replace(/\.[^.]+$/, "").replace(/[-_]/g, " ");
|
||
const parts = [`${file} — ${sizeKb}KB`];
|
||
if (geminiCaption) {
|
||
parts.push(geminiCaption);
|
||
captionedLines.push(parts.join(", "));
|
||
} else {
|
||
if (desc) parts.push(`"${desc.slice(0, 80)}"`);
|
||
if (heading) parts.push(`section: "${heading.slice(0, 60)}"`);
|
||
else if (section) parts.push(`in: ${section.split(" ").slice(0, 3).join(" ")}`);
|
||
if (aboveFold) parts.push(aboveFold);
|
||
if (!desc && !heading) parts.push(cleanName);
|
||
uncaptionedLines.push(parts.join(", "));
|
||
}
|
||
}
|
||
} catch {
|
||
/* no assets dir */
|
||
}
|
||
|
||
// Describe SVGs
|
||
try {
|
||
const svgsPath = join(assetsPath, "svgs");
|
||
for (const file of readdirSync(svgsPath)) {
|
||
if (!file.endsWith(".svg")) continue;
|
||
const svgMatch = tokens.svgs.find(
|
||
(s) =>
|
||
s.label &&
|
||
file.includes(
|
||
s.label
|
||
.toLowerCase()
|
||
.replace(/[^a-z0-9]/g, "-")
|
||
.slice(0, 15),
|
||
),
|
||
);
|
||
const geminiCaption = geminiCaptions[`svgs/${file}`];
|
||
if (geminiCaption) {
|
||
svgLines.push(`svgs/${file} — ${geminiCaption}`);
|
||
continue;
|
||
}
|
||
const label = svgMatch?.label || file.replace(".svg", "").replace(/-/g, " ");
|
||
svgLines.push(`svgs/${file} — ${label}`);
|
||
}
|
||
} catch {
|
||
/* no svgs dir */
|
||
}
|
||
|
||
// Describe fonts
|
||
try {
|
||
const fontsPath = join(assetsPath, "fonts");
|
||
for (const file of readdirSync(fontsPath)) {
|
||
fontLines.push(`fonts/${file} — font file`);
|
||
}
|
||
} catch {
|
||
/* no fonts dir */
|
||
}
|
||
|
||
// Describe videos — high-value motion clips. The video-manifest.json (written
|
||
// earlier by captureVideoManifest) carries each clip's DOM heading/caption +
|
||
// dims. Surfaced FIRST and tagged `[video]`: for a product/demo these moving
|
||
// clips are usually the strongest hero material, and downstream planners key off
|
||
// the `[video]` marker. (The `videos/` dir is skipped in the image walk above —
|
||
// its entries come from the manifest, which has the captions the bare files lack.)
|
||
const videoLines: string[] = [];
|
||
try {
|
||
const manifest = JSON.parse(
|
||
readFileSync(join(outputDir, "extracted", "video-manifest.json"), "utf-8"),
|
||
) as Array<{
|
||
filename?: string;
|
||
localPath?: string;
|
||
caption?: string;
|
||
heading?: string;
|
||
width?: number;
|
||
height?: number;
|
||
sourceWidth?: number;
|
||
sourceHeight?: number;
|
||
}>;
|
||
for (const v of manifest) {
|
||
if (!v.localPath) continue; // only describe clips that actually downloaded
|
||
const base = basename(v.localPath) || v.filename || "";
|
||
if (!base) continue;
|
||
const desc =
|
||
(v.caption || v.heading || "").trim().replace(/\s+/g, " ").slice(0, 140) || "motion clip";
|
||
const dimW = v.sourceWidth || v.width;
|
||
const dimH = v.sourceHeight || v.height;
|
||
const dims = dimW && dimH ? `, ~${dimW}×${dimH}` : "";
|
||
videoLines.push(`${base} — [video] ${desc}${dims}`);
|
||
}
|
||
} catch {
|
||
/* no video manifest */
|
||
}
|
||
|
||
return [...videoLines, ...captionedLines, ...uncaptionedLines, ...svgLines, ...fontLines];
|
||
}
|