Implementing DeepSeek V4

This commit is contained in:
Owen Qwen
2026-08-05 21:13:25 -05:00
parent dbdf0de9b7
commit cabe22c4c0
9 changed files with 109 additions and 26 deletions
+1 -1
View File
@@ -2631,7 +2631,7 @@ common_params_context common_params_parser_init(common_params & params, llama_ex
).set_env("LLAMA_ARG_LOAD_MODE"));
add_opt(common_arg(
{"--longhaul"},
"stream routed Qwen3.5 MoE, Laguna, or Inkling experts through a bounded CPU or Metal cache",
"stream routed Qwen3.5 MoE, DeepSeek V4, Laguna, or Inkling experts through a bounded CPU or Metal cache",
[](common_params & params) {
params.load_mode = LLAMA_LOAD_MODE_LONGHAUL;
}