]> git.djapps.eu Git - pkg/ggml/sources/llama.cpp/commitdiff
ggml-hip: enable -funsafe-math-optimizations (#24668)
authorRapidMark <redacted>
Thu, 9 Jul 2026 08:02:26 +0000 (01:02 -0700)
committerGitHub <redacted>
Thu, 9 Jul 2026 08:02:26 +0000 (11:02 +0300)
CUDA is compiled with fast math and AMD/HIP is not — this flag lets AMD use fast math too.

We can't use -ffast-math: it implies -ffinite-math-only, which won't compile (ggml uses INFINITY for masking) and produces NaNs. -funsafe-math-optimizations gives the speedup without the NaN problems.

Co-authored-by: Mark Caldwell <redacted>
ggml/src/ggml-hip/CMakeLists.txt

index 44fd3f3c5bef2c7380cd38fa0dfdf9b152073947..7121193f1c84e9ba5b694ae3b74597534e0192fc 100644 (file)
@@ -130,6 +130,9 @@ if (GGML_HIP_EXPORT_METRICS)
     set(CMAKE_HIP_FLAGS "${CMAKE_HIP_FLAGS} -Rpass-analysis=kernel-resource-usage --save-temps")
 endif()
 
+# Fast math for HIP, like CUDA's -use_fast_math. Not -ffast-math: that implies -ffinite-math-only, which breaks ggml's INFINITY masking and produces NaNs.
+set(CMAKE_HIP_FLAGS "${CMAKE_HIP_FLAGS} -funsafe-math-optimizations")
+
 if (NOT GGML_CUDA_FA)
     add_compile_definitions(GGML_CUDA_NO_FA)
 endif()