yash184 commited on
Commit
ec50f8d
·
verified ·
1 Parent(s): defbe9e

Update Dockerfile

Browse files
Files changed (1) hide show
  1. Dockerfile +37 -26
Dockerfile CHANGED
@@ -1,29 +1,40 @@
1
- # Base image llama.cpp se hi le rahe hain
2
- FROM ghcr.io/ggml-org/llama.cpp:full
 
 
 
 
 
 
 
 
 
 
 
3
 
4
  WORKDIR /app
5
 
6
- # System dependencies install karna
7
- RUN apt update && apt install -y python3 python3-pip python3-venv
8
-
9
- # Virtual environment setup
10
- RUN python3 -m venv /opt/venv
11
- ENV PATH="/opt/venv/bin:$PATH"
12
-
13
- # Python tools install karna
14
- RUN pip install -U pip huggingface_hub
15
-
16
- # Model download karna (Qwen2.5-Coder-1.5B)
17
- RUN python3 -c 'from huggingface_hub import hf_hub_download; \
18
- repo="bartowski/Qwen2.5-Coder-1.5B-Instruct-GGUF"; \
19
- hf_hub_download(repo_id=repo, filename="Qwen2.5-Coder-1.5B-Instruct-Q6_K.gguf", local_dir="/app")'
20
-
21
- # Server start karne ka command
22
- # Humne mmproj hata diya hai taaki error na aaye
23
- CMD [
24
- "llama-server",
25
- "-m", "/app/Qwen2.5-Coder-1.5B-Instruct-Q6_K.gguf",
26
- "--host", "0.0.0.0",
27
- "--port", "7860",
28
- "-c", "32768",
29
- "-ngl", "0"]
 
1
+ FROM ubuntu:24.04
2
+
3
+ ENV DEBIAN_FRONTEND=noninteractive
4
+
5
+ RUN apt-get update && \
6
+ apt-get install -y \
7
+ git \
8
+ build-essential \
9
+ cmake \
10
+ python3 \
11
+ python3-pip \
12
+ curl && \
13
+ rm -rf /var/lib/apt/lists/*
14
 
15
  WORKDIR /app
16
 
17
+ # Build latest llama.cpp
18
+ RUN git clone https://github.com/ggml-org/llama.cpp.git && \
19
+ cd llama.cpp && \
20
+ cmake -B build && \
21
+ cmake --build build -j$(nproc)
22
+
23
+ RUN pip3 install --no-cache-dir huggingface_hub
24
+
25
+ RUN python3 -c "\
26
+ from huggingface_hub import hf_hub_download; \
27
+ hf_hub_download( \
28
+ repo_id='bartowski/Qwen2.5-Coder-1.5B-Instruct-GGUF', \
29
+ filename='Qwen2.5-Coder-1.5B-Instruct-Q6_K.gguf', \
30
+ local_dir='/app/model' \
31
+ )"
32
+
33
+ EXPOSE 7860
34
+
35
+ CMD ["/app/llama.cpp/build/bin/llama-server",
36
+ "-m","/app/model/Qwen2.5-Coder-1.5B-Instruct-Q6_K.gguf",
37
+ "--host","0.0.0.0",
38
+ "--port","7860",
39
+ "-c","32768",
40
+ "-ngl","0"]