fix: QA sweep — 7 bugs fixed, 17 test corrections

Code fixes:
- Sidebar state bleed: reset file store on HomePage mount
- restore-photo: raise error instead of silently skipping colorize
  when DDColor model missing
- PaddleOCR OOM: cap input images to 2048px before OCR inference
- Torch CPU optimization: use --index-url .../whl/cpu on CPU nodes

Test fixes:
- upscale: add exact:true to scale factor button locators
- smart-crop: add exact:true to "Pad to square" locator
- colorize: use regex for model button names (Best/Balanced/Fast)
- enhance-faces: use .first() for ambiguous percentage display
- passport-photo: fix DPI locator, .or() compound, generate fallback
- people: update maxUsers assertions for unlimited (0) default
- automate: "Save Pipeline" → "Save" matching actual button text
- tools.test: add resize to Sharp mock chain for OCR tests
This commit is contained in:
SnapOtter
2026-04-25 07:23:58 +08:00
parent bd84728588
commit bf0307d87d
12 changed files with 47 additions and 50 deletions
+4 -2
View File
@@ -75,7 +75,7 @@ def cpu_fallback_packages(packages: list[str]) -> list[str]:
# "torch==2.7.0+cu126 torchvision==0.22.0+cu126 --index-url ..."
first_token = pkg.split()[0] if pkg.strip() else ""
if first_token.startswith("torch==") and "+cu" in first_token:
# Extract torch and torchvision versions, strip CUDA suffix
# Extract torch and torchvision versions, use CPU-only index
cpu_pkgs = []
for token in pkg.split():
if token.startswith("torch==") and "+cu" in token:
@@ -84,7 +84,9 @@ def cpu_fallback_packages(packages: list[str]) -> list[str]:
elif token.startswith("torchvision==") and "+cu" in token:
base_ver = token.split("+")[0] # "torchvision==0.21.0"
cpu_pkgs.append(base_ver)
# Drop --index-url and its argument (not needed for CPU torch)
# Use CPU-only wheels (~200MB vs ~2.6GB with CUDA)
cpu_pkgs.append("--index-url")
cpu_pkgs.append("https://download.pytorch.org/whl/cpu")
result.extend(cpu_pkgs)
continue
+4 -1
View File
@@ -483,7 +483,10 @@ def colorize_bw(img_bgr, intensity=0.85):
from gpu import safe_onnx_session
if not os.path.exists(DDCOLOR_MODEL_PATH):
return img_bgr, False
raise FileNotFoundError(
f"DDColor model not found at {DDCOLOR_MODEL_PATH}. "
"Install the 'object-eraser-colorize' bundle to enable colorization."
)
session, _device = safe_onnx_session(DDCOLOR_MODEL_PATH)
input_name = session.get_inputs()[0].name
+7 -4
View File
@@ -26,12 +26,15 @@ export async function extractText(
): Promise<OcrResult> {
const inputPath = join(outputDir, "input_ocr.png");
// Convert any input format (HEIC, AVIF, WebP, TIFF, etc.) to PNG
// so Tesseract and PaddleOCR can read it reliably.
const pngBuffer = await sharp(inputBuffer).png().toBuffer();
// Convert to PNG and cap at 2048px to prevent PaddleOCR OOM on large images.
const MAX_OCR_DIM = 2048;
const pngBuffer = await sharp(inputBuffer)
.resize({ width: MAX_OCR_DIM, height: MAX_OCR_DIM, fit: "inside", withoutEnlargement: true })
.png()
.toBuffer();
await writeFile(inputPath, pngBuffer);
const meta = await sharp(inputBuffer).metadata();
const meta = await sharp(pngBuffer).metadata();
const megapixels = ((meta.width ?? 0) * (meta.height ?? 0)) / 1_000_000;
const timeout = Math.max(600_000, megapixels * 30 * 1000);