]> git.djapps.eu Git - pkg/ggml/sources/whisper.cpp/commitdiff
cuda: fix rope fusion for gemma3 (llama/17378)
authorAman Gupta <redacted>
Wed, 19 Nov 2025 10:25:05 +0000 (18:25 +0800)
committerGeorgi Gerganov <redacted>
Fri, 12 Dec 2025 15:53:03 +0000 (17:53 +0200)
ggml/src/ggml-cuda/ggml-cuda.cu

index 7d792e60cf9c5e32c74b521ec6c67184f646825e..889801cb5dadf603242ebddcb9da9042398e231e 100644 (file)
@@ -3001,6 +3001,10 @@ static void update_cuda_graph_executable(ggml_backend_cuda_context * cuda_ctx) {
 static bool ggml_cuda_should_fuse_rope_set_rows(const ggml_tensor * rope,
                                                 const ggml_tensor * view,
                                                 const ggml_tensor * set_rows) {
+
+    if (rope->op != GGML_OP_ROPE || view->op != GGML_OP_VIEW || set_rows->op != GGML_OP_SET_ROWS) {
+        return false;
+    }
     // ne3 not tested
     if (rope->src[0]->ne[3] != 1) {
         return false;