Merge pull request #3 from cardinalfan1/claude/fix-runpod-startup-xxJqT
Fix -m None passed to llama-server when cached model not found
This commit is contained in:
@@ -3,6 +3,7 @@ Finds the full LLM GGUF path from the Hugging Face cache.
|
||||
"""
|
||||
|
||||
import os
|
||||
import sys
|
||||
import argparse
|
||||
|
||||
CACHE_DIR = "/runpod-volume/huggingface-cache/hub"
|
||||
@@ -52,6 +53,12 @@ def main():
|
||||
args = parser.parse_args()
|
||||
|
||||
model_path = find_model_path(args.model, args.path)
|
||||
if model_path is None:
|
||||
print(
|
||||
f"Error: Cached model not found. Model='{args.model}', GGUF='{args.path}', Cache dir='{CACHE_DIR}'",
|
||||
file=sys.stderr,
|
||||
)
|
||||
sys.exit(1)
|
||||
print(model_path, end="")
|
||||
|
||||
|
||||
|
||||
+7
-1
@@ -17,7 +17,13 @@ cleanup() {
|
||||
CACHED_LLAMA_ARGS=""
|
||||
|
||||
find_cached_path() {
|
||||
CACHED_LLAMA_ARGS="-m $(python ./find_cached.py $LLAMA_CACHED_MODEL $LLAMA_CACHED_GGUF_PATH)"
|
||||
local model_path
|
||||
model_path=$(python ./find_cached.py "$LLAMA_CACHED_MODEL" "$LLAMA_CACHED_GGUF_PATH")
|
||||
if [ $? -ne 0 ] || [ -z "$model_path" ]; then
|
||||
echo "start.sh: Error: Could not resolve cached model path. Check that LLAMA_CACHED_MODEL and LLAMA_CACHED_GGUF_PATH are correct and the model is fully cached on the network volume."
|
||||
exit 1
|
||||
fi
|
||||
CACHED_LLAMA_ARGS="-m $model_path"
|
||||
}
|
||||
|
||||
# check if $LLAMA_CACHED_MODEL is set and not empty
|
||||
|
||||
Reference in New Issue
Block a user