mirror of
https://github.com/ggml-org/llama.cpp.git
synced 2026-10-11 07:20:33 +02:00
Avoid races in fp8_fallback
This commit is contained in:
@@ -24,6 +24,7 @@ static __global__ void mul_mat_fp8_fallback(
|
||||
if (threadIdx.x == 0) {
|
||||
*(float *) (dst + i0*sizeof(float) + i1*nb1 + i2*nb2 + i3*nb3) = sum;
|
||||
}
|
||||
__syncthreads();
|
||||
}
|
||||
}
|
||||
|
||||
|
||||
Reference in New Issue
Block a user