Merge branch 'master' into gg/llama-kv-cache

ggml-ci
2025-11-16 11:27:03 +00:00 · 2025-02-18 10:14:37 +02:00
parent 1d801d27b9 73e2ed3ce3
commit f0d3ff2388
156 changed files with 6433 additions and 2603 deletions
--- a/src/llama-kv-cache.h
+++ b/src/llama-kv-cache.h
@@ -57,7 +57,7 @@ struct llama_kv_cache {
    bool can_shift = false;

    // Note: The value of head isn't only used to optimize searching
-    // for a free KV slot. llama_decode_internal also uses it, so it
+    // for a free KV slot. llama_decode_impl also uses it, so it
    // cannot be freely changed after a slot has been allocated.
    uint32_t head = 0;
    uint32_t size = 0;