# FabledCurator GPU agent — runs on the desktop with the GPU. # CUDA 12.9 + cuDNN 9 runtime so onnxruntime-gpu can use the card (it needs # cuDNN 9 — the plain -runtime image lacks it: "libcudnn.so.9: cannot open # shared object file"); ffmpeg for video frames. Ubuntu 24.04 → Python 3.12. # Stays on the CUDA-12 / cuDNN-9 line the default onnxruntime-gpu + torch are # built against (CUDA 13 has only nascent ONNX Runtime support). FROM nvidia/cuda:12.9.2-cudnn-runtime-ubuntu24.04 # PIP_BREAK_SYSTEM_PACKAGES: Ubuntu 24.04 marks its system Python as externally # managed (PEP 668), so a global `pip install` errors without this. It's a # single-purpose container — we own the whole environment, so installing into # the system site-packages is fine (and simplest — no venv on PATH to manage). ENV DEBIAN_FRONTEND=noninteractive PYTHONUNBUFFERED=1 PIP_BREAK_SYSTEM_PACKAGES=1 RUN apt-get update \ && apt-get install -y --no-install-recommends python3 python3-pip ffmpeg \ && rm -rf /var/lib/apt/lists/* WORKDIR /app # torch from the CUDA-12.4 wheel index; its wheels bundle their own CUDA + cuDNN # so they run on the 12.9 base and coexist with onnxruntime-gpu. Installed first # + separately so the GPU build of torch is deterministic and layer-cached. RUN pip3 install --no-cache-dir torch==2.6.0 --index-url https://download.pytorch.org/whl/cu124 COPY requirements.txt . RUN pip3 install --no-cache-dir -r requirements.txt COPY fc_agent ./fc_agent # imgutils ONNX models + the transformers SigLIP weights both cache here; mount # a volume to persist them across restarts (the SigLIP download is ~3.5 GB once). ENV HF_HOME=/models # Declared LAST on purpose, exactly as the web Dockerfile does: an ARG/ENV # invalidates every layer below it, and these are the only values that differ # between builds of otherwise identical source. Any earlier and the ~6.3 GB # CUDA + torch layers could never be shared between the dev and main builds of # one commit — which is the cost #3114 measured at 9m26s cold. # # Three values, never folded together (rule 149) — the NAME a person reads, the # CHANNEL it came from, and the REVISION that identifies the content. See # fc_agent/build_info.py; CI derives all three from scripts/artifacts.sh. ARG FC_CHANNEL="" ENV FC_CHANNEL=${FC_CHANNEL} ARG FC_VERSION="" ENV FC_VERSION=${FC_VERSION} ARG FC_REVISION="" ENV FC_REVISION=${FC_REVISION} EXPOSE 8770 # The control UI; the worker is started from it (or POST /start). CMD ["uvicorn", "fc_agent.app:app", "--host", "0.0.0.0", "--port", "8770"]