bench-server / Dockerfile
tinkersnot's picture
Upload folder using huggingface_hub
b20f4bd verified
Raw
History Blame Contribute Delete
688 Bytes
# tinkersnot bake-off inference Space — llama.cpp server, own entrypoint for full control.
FROM ghcr.io/ggml-org/llama.cpp:server
# HF Spaces injects a `git config` build step; the minimal ggml image lacks git.
USER root
RUN apt-get update && apt-get install -y --no-install-recommends git && rm -rf /var/lib/apt/lists/*
# Bake the GGUF at build time (no runtime download / hang). SMOKE TEST: Qwen3-0.6B.
ADD https://huggingface.co/ggml-org/Qwen3-0.6B-GGUF/resolve/main/Qwen3-0.6B-Q8_0.gguf /models/model.gguf
# Override the base image entrypoint entirely — our script prints boot info then execs the server.
COPY start.sh /start.sh
RUN chmod +x /start.sh
ENTRYPOINT ["/start.sh"]