Commit a55e952b8 for llama.cpp
commit a55e952b85741249fcbf120bcd9d47cf2b1f6e48
Author: Pascal <admin@serveurperso.com>
Date: Sat Oct 3 14:56:07 2026 +0200
ci: fix flaky ADD_ADD f16 by using the fused ADD tolerance (#29904)
diff --git a/tests/test-backend-ops.cpp b/tests/test-backend-ops.cpp
index 5c7f61a78..37d3ef086 100644
--- a/tests/test-backend-ops.cpp
+++ b/tests/test-backend-ops.cpp
@@ -3913,6 +3913,17 @@ struct test_add_add : public test_case {
return out;
}
+
+ double max_nmse_err() override {
+ // Fused ADDs can keep FP32 intermediates while the CPU rounds each ADD to FP16/BF16.
+ if (type == GGML_TYPE_F16) {
+ return 1e-6;
+ }
+ if (type == GGML_TYPE_BF16) {
+ return 1e-4;
+ }
+ return test_case::max_nmse_err();
+ }
};
// GGML_OP_ADD + GGML_OP_RMS_NORM (fused operation)