2 Commits
Author SHA1 Message Date
mags0ft 381ed9ffff prepare for production release 2025-11-19 18:40:00 +01:00
mags0ft c42b80ebbd add default behavior for undefined LLAMA_SERVER_CMD_ARGS 2025-11-19 18:11:36 +01:00
3 changed files with 7 additions and 2 deletions
+1 -1
View File
@@ -14,7 +14,7 @@
{
"key": "LLAMA_SERVER_CMD_ARGS",
"input": {
"name": "Model Name",
"name": "Command line arguments for llama-server",
"type": "string",
"description": "Launch command line arguments (argv) for the llama-server binary. Do not define the port.",
"default": "-hf unsloth/gemma-3-270m-it-GGUF:Q6_K --ctx-size 4096",
-1
View File
@@ -21,7 +21,6 @@ Typical usage:
"""
import json
import os
from dotenv import load_dotenv
from openai import OpenAI
+6
View File
@@ -14,6 +14,12 @@ cleanup() {
exit 0
}
# check if $LLAMA_SERVER_CMD_ARGS is set
if [ -z "$LLAMA_SERVER_CMD_ARGS" ]; then
echo "start.sh: Warning: LLAMA_SERVER_CMD_ARGS is not set. Defaulting to -hf unsloth/gemma-3-270m-it-GGUF:Q6_K --ctx-size 4096"
LLAMA_SERVER_CMD_ARGS="-hf unsloth/gemma-3-270m-it-GGUF:Q6_K --ctx-size 4096"
fi
# check if the substring /workspace is in LLAMA_SERVER_CMD_ARGS
if [[ "$LLAMA_SERVER_CMD_ARGS" != *"/workspace"* ]]; then
echo "start.sh: Tip: For reduced downloads and faster startup times, consider using a model stored in a network volume mounted to /workspace."