diff --git a/Dockerfile b/Dockerfile index 2bd0c27..2ae902c 100644 --- a/Dockerfile +++ b/Dockerfile @@ -21,10 +21,10 @@ RUN --mount=type=cache,target=/root/.cache/pip \ # Install torch and vllm based on CUDA version RUN if [[ "${WORKER_CUDA_VERSION}" == 11.8* ]]; then \ - python3.11 -m pip install -e git+https://github.com/alpayariyak/vllm.git@cuda-11.8#egg=vllm; \ + python3.11 -m pip install -e git+https://github.com/runpod/vllm-fork-for-sls-worker.git@cuda-11.8#egg=vllm; \ python3.11 -m pip install -U --force-reinstall torch==2.1.2 xformers==0.0.23.post1 --index-url https://download.pytorch.org/whl/cu118; \ else \ - python3.11 -m pip install -e git+https://github.com/alpayariyak/vllm.git#egg=vllm; \ + python3.11 -m pip install -e git+https://github.com/runpod/vllm-fork-for-sls-worker.git#egg=vllm; \ fi && \ rm -rf /root/.cache/pip @@ -47,4 +47,4 @@ RUN if [ -n "$MODEL_NAME" ]; then \ fi # Start the handler -CMD ["python3.11", "/handler.py"] \ No newline at end of file +CMD ["python3.11", "/handler.py"]