whisper.cpp / ggml /src /ggml-cuda /template-instances /fattn-mma-f16-instance-ncols1_8-ncols2_2.cu
JohannesGaessler's picture
CUDA: FA support for Deepseek (Ampere or newer) (llama/13306)
507d30c
// This file has been autogenerated by generate_cu_files.py, do not edit manually.
#include "../fattn-mma-f16.cuh"
DECL_FATTN_MMA_F16_CASE(64, 64, 8, 2);
DECL_FATTN_MMA_F16_CASE(80, 80, 8, 2);
DECL_FATTN_MMA_F16_CASE(96, 96, 8, 2);
DECL_FATTN_MMA_F16_CASE(112, 112, 8, 2);
DECL_FATTN_MMA_F16_CASE(128, 128, 8, 2);
DECL_FATTN_MMA_F16_CASE(256, 256, 8, 2);