c9384636e625f1b10868552d8add22220d6f5b33

Author
TheEdgeOfRage <git@theedgeofrage.com>
Committer
TheEdgeOfRage <git@theedgeofrage.com>
Date

Message

Add supergemma4 model to llama-server

Diff

 1diff --git a/dot_config/llama-server/models.ini b/dot_config/llama-server/models.ini
 2index 693b839ade6a759141b84d1be5661fb3f873dc53..b26aff5f62069e0ebd8dccf320e5ea3fc05d8b78 100644
 3--- a/dot_config/llama-server/models.ini
 4+++ b/dot_config/llama-server/models.ini
 5@@ -34,6 +34,18 @@ spec-draft-n-max = 2
 6 repeat-penalty = 1.0
 7 reasoning = on
 8 
 9+[supergemma4-26b]
10+hf = Jiunsong/supergemma4-26b-uncensored-gguf-v2:Q4_K_M
11+alias = supergemma4-26b
12+n-gpu-layers = 99
13+flash-attn = on
14+ctx-size = 65536
15+temp = 1.1
16+top-p = 0.95
17+top-k = 64
18+fit = off
19+reasoning = auto
20+
21 # [gemma4-26b]
22 # hf = unsloth/gemma-4-26B-A4B-it-GGUF:UD-Q4_K_M
23 # hf-draft = unsloth/gemma-4-26B-A4B-it-GGUF:Q8_0-MTP